<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>aictrl.dev Blog</title>
    <link>https://aictrl.dev/blog</link>
    <description>Research, analysis, and practical guides for engineering teams deploying AI agents in production.</description>
    <language>en</language>
    <lastBuildDate>Thu, 01 Oct 2026 00:00:00 GMT</lastBuildDate>
    <atom:link href="https://aictrl.dev/blog/feed.xml" rel="self" type="application/rss+xml"/>
    <item>
      <title>The AI Software Factory: Managing Change for Measurable Impact</title>
      <link>https://aictrl.dev/blog/software-factory-change-impact</link>
      <description>An AI software factory connects managed inputs, controlled change, and measured impact. A practical model for workflows, integration, economics, and Product strategy.</description>
      <pubDate>Thu, 01 Oct 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/software-factory-change-impact</guid>
    </item>
    <item>
      <title>We simplified our backlog screen, and users would have approved the wrong thing | aictrl.dev</title>
      <link>https://aictrl.dev/blog/ai-usability-testing-backlog-design</link>
      <description>We cut our Backlog screen from 37 controls to 16. Then AI novice testers approved a security review instead of starting work. What our tests caught.</description>
      <pubDate>Sat, 26 Sep 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/ai-usability-testing-backlog-design</guid>
    </item>
    <item>
      <title>From a generated image to a Blender film | aictrl.dev</title>
      <link>https://aictrl.dev/blog/from-agent-factory-to-workflow</link>
      <description>How we made the aictrl homepage film: a generated reference, an authored 3D world, and a workflow transition—with the failures and fixes along the way.</description>
      <pubDate>Wed, 09 Sep 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/from-agent-factory-to-workflow</guid>
    </item>
    <item>
      <title>Get It Off the Laptop: From Interactive AI Loops to Measured Workflows</title>
      <link>https://aictrl.dev/blog/agentic-loops-to-measured-workflows</link>
      <description>Interactive AI harnesses and automated workflows are different machines, not two settings of one dial. Attribution requires boundaries — a loop that leaves no evidence never compounds.</description>
      <pubDate>Wed, 22 Jul 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/agentic-loops-to-measured-workflows</guid>
    </item>
    <item>
      <title>AI Didn't Kill the SDLC. It Compressed It.</title>
      <link>https://aictrl.dev/blog/ai-sdlc-compressed-not-replaced</link>
      <description>1,456 merged PRs and 155 production deploys in 90 days from a one-engineer team. The SDLC phases survived — compressed into AI loops with human gates, measured by DORA metrics.</description>
      <pubDate>Fri, 18 Jul 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/ai-sdlc-compressed-not-replaced</guid>
    </item>
    <item>
      <title>Code Review With a 12B Model: Graph Topology and the Price of Recall</title>
      <link>https://aictrl.dev/blog/local-model-code-review</link>
      <description>A recall-first DAG pipeline around gemma-12B found real TypeScript review bugs by combining many noisy passes, graph topologies, and script checks.</description>
      <pubDate>Sun, 07 Jun 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/local-model-code-review</guid>
    </item>
    <item>
      <title>aictrl vs. GitHub vs. GitLab for Agentic SDLC Automation</title>
      <link>https://aictrl.dev/blog/aictrl-vs-github-gitlab-sdlc-automation</link>
      <description>A decision-maker guide to where GitHub, GitLab, and aictrl.dev fit across issue-to-PR automation, DevSecOps, governed skills, triggered workflows, observability, and cross-tool orchestration.</description>
      <pubDate>Mon, 01 Jun 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/aictrl-vs-github-gitlab-sdlc-automation</guid>
    </item>
    <item>
      <title>How We Evaluated AI Code Review on Our Own PRs</title>
      <link>https://aictrl.dev/blog/code-review-dogfooding-assessment-2026</link>
      <description>A dogfooding assessment over 212 anonymized pull requests showing why valid finding rate was not enough and why material and weighted FIX coverage changed the primary-reviewer decision.</description>
      <pubDate>Thu, 28 May 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/code-review-dogfooding-assessment-2026</guid>
    </item>
    <item>
      <title>AI Code Review Tools in 2026: The PR Surface Is the Product</title>
      <link>https://aictrl.dev/blog/ai-code-review-tools-2026</link>
      <description>A practical review of CodeRabbit, CodeAnt, GitHub Copilot Code Review, Qodo, Graphite, and Cursor Bugbot across PR features, inline comments, analytics, team training, docs, and eval credibility.</description>
      <pubDate>Mon, 25 May 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/ai-code-review-tools-2026</guid>
    </item>
    <item>
      <title>Why Knowledge Graphs Are the Missing Infrastructure Layer for Agentic AI</title>
      <link>https://aictrl.dev/blog/knowledge-graph-agentic-ai</link>
      <description>95% of enterprise AI pilots fail. The root cause isn&apos;t the models — it&apos;s the data layer. Research shows knowledge graphs improve LLM accuracy by 3-5x and cut token costs by 80-97%.</description>
      <pubDate>Thu, 26 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/knowledge-graph-agentic-ai</guid>
    </item>
    <item>
      <title>From Vibe Coding to Spec Engineering: The Discipline Gap Between Reliable and Unpredictable AI Teams</title>
      <link>https://aictrl.dev/blog/skill-engineering-pseudocode</link>
      <description>Agent instruction quality drives a 2.3x performance swing — more than model choice. Research shows why top AI teams treat instructions as production code.</description>
      <pubDate>Thu, 19 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/skill-engineering-pseudocode</guid>
    </item>
    <item>
      <title>The SKILL.md Trajectory: From Open Standard to Industry Default in 90 Days</title>
      <link>https://aictrl.dev/blog/skill-md-adoption-trajectory</link>
      <description>Anthropic open-sourced Agent Skills on Dec 18. OpenAI adopted within 48 hours. By February, 160,000+ skills indexed across 7+ platforms. Adoption map and infrastructure forecast.</description>
      <pubDate>Wed, 11 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/skill-md-adoption-trajectory</guid>
    </item>
    <item>
      <title>The SKILL.md Standard: Your Enterprise Guide to AI Grounding</title>
      <link>https://aictrl.dev/blog/skill-md-standard-enterprise-guide</link>
      <description>84% of developers use AI tools, but only 23% of enterprises can measure ROI. How structured instruction files deliver 20-50% performance gains with a step-by-step plan.</description>
      <pubDate>Sat, 07 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/skill-md-standard-enterprise-guide</guid>
    </item>
    <item>
      <title>The Reasoning Race: From 2.7% to 53.1% on Humanity&apos;s Last Exam</title>
      <link>https://aictrl.dev/blog/reasoning-benchmarks</link>
      <description>AI reasoning exploded in 12 months. Claude Opus 4.6 leads HLE at 53.1%, GPT-5.2 dominates math. Analysis of benchmark saturation, model specialization, and enterprise ROI.</description>
      <pubDate>Thu, 05 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/reasoning-benchmarks</guid>
    </item>
    <item>
      <title>The SaaSpocalypse: $1 Trillion Wipeout and Who Survives</title>
      <link>https://aictrl.dev/blog/saaspocalypse</link>
      <description>Anthropic&apos;s Claude Cowork triggered a ~$1 trillion selloff in software stocks. Legal tech fell 20%, SaaS giants lost 26-59% YTD. Analysis of who survived and why.</description>
      <pubDate>Thu, 05 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/saaspocalypse</guid>
    </item>
    <item>
      <title>The Swarm Paradox: Why More AI Agents Often Means Worse Results</title>
      <link>https://aictrl.dev/blog/swarm-paradox</link>
      <description>Google found multi-agent systems degrade sequential tasks by 39-70%. Agents amplify errors 17.2x. Analysis of 12 benchmarks and what actually works.</description>
      <pubDate>Thu, 05 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/swarm-paradox</guid>
    </item>
    <item>
      <title>The AI Electricity Crisis: Will Efficiency Save Us?</title>
      <link>https://aictrl.dev/blog/ai-electricity-crisis</link>
      <description>GPU efficiency improved 66x in 9 years. Token demand grew 3,000x. Analysis of Jevons Paradox in AI, state-level electricity data, and what CTOs should do about the growing gap.</description>
      <pubDate>Wed, 04 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/ai-electricity-crisis</guid>
    </item>
    <item>
      <title>The Rise of AI-Native Browser Automation</title>
      <link>https://aictrl.dev/blog/browser-use-rise</link>
      <description>How browser-use went from 0 to 77,000 GitHub stars in 5 months, challenging a decade of testing infrastructure. Analysis of the shift from Playwright to AI-native automation.</description>
      <pubDate>Tue, 03 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/browser-use-rise</guid>
    </item>
    <item>
      <title>Framework Matters: The 2.5x AI Performance Story</title>
      <link>https://aictrl.dev/blog/framework-matters</link>
      <description>Analysis of how AI coding agents perform differently across frameworks. Data showing why your choice of tools significantly impacts AI-assisted development outcomes.</description>
      <pubDate>Mon, 02 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/framework-matters</guid>
    </item>
    <item>
      <title>SWE-Bench Evolution Analysis: AI Coding Agent Performance</title>
      <link>https://aictrl.dev/blog/swe-evo-analysis</link>
      <description>Deep dive into how AI coding benchmarks have evolved and what the latest results tell us about the state of autonomous software engineering.</description>
      <pubDate>Mon, 02 Feb 2026 00:00:00 GMT</pubDate>
      <guid isPermaLink="true">https://aictrl.dev/blog/swe-evo-analysis</guid>
    </item>
  </channel>
</rss>
