{"id":"W4416873480","doi":"10.1109/inc465408.2025.11256489","title":"Enhancing Text-to-Image Generation Using Ensemble Vision Transformers","year":2025,"lang":"","type":"article","venue":"","topic":"Generative Adversarial Networks and Image Synthesis","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Optech (Canada)","funders":"","keywords":"Bridging (networking); Transformer; Modalities; Consistency (knowledge bases); Metric (unit); Representation (politics)","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001043112,0.001012882,0.000640791,0.0005660229,0.0002076058,0.0008325299,0.001283447,0.0008341109,0.004117388],"category_scores_gemma":[0.003472863,0.0002831379,0.000663778,0.0003359543,0.0005374224,0.001450888,0.001302268,0.001475585,0.001213687],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005406824,"about_ca_system_score_gemma":0.0004588893,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001508483,"about_ca_topic_score_gemma":0.002044557,"domain_scores_codex":[0.9995375,0.0001136316,0.00001708113,0.0001334437,0.0001495463,0.00004887065],"domain_scores_gemma":[0.998861,0.0006554579,0.00005565672,0.0001907126,0.0001749551,0.0000624021],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0005665457,0.0003750781,0.001083387,0.0002010267,0.00009419218,0.0002986333,0.0001558624,0.4524984,0.06591685,0.01007441,0.00818285,0.4605528],"study_design_scores_gemma":[0.00001731618,0.00007092817,0.0001100816,0.00000487289,0.000008728578,0.00005684534,0.00001118168,0.9835127,0.01252529,0.002799337,0.0008759801,0.000006738615],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.03700788,0.0003737659,0.9509482,0.0002045314,0.0001501111,0.0001190233,0.0001647773,0.007037037,0.003994612],"genre_scores_gemma":[0.7414734,0.0002774225,0.2489347,0.0003859081,0.00009676973,0.0001299344,0.0009539788,0.0007965707,0.006951323],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.004117388,"threshold_uncertainty_score":0.0137741,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02009932331468428,"score_gpt":0.284973804288421,"score_spread":0.2648744809737367,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}