test(#1074): add additive per-file workflow size baseline guard (PR 1/3) (#1089)

* test(#1074): add additive per-file workflow size baseline guard (PR 1/3)

Introduces a committed per-file size baseline scheme alongside (not replacing)
the existing tier anti-creep tests. Green by construction — the baseline
records current sizes, so both schemes pass side by side during migration.

- scripts/lib/allowlist-ratchet.cjs: add assertFileBaseline (third pure helper,
  same injected-fail style) — per-file growth/shrink/add/remove diff vs baseline.
- scripts/workflow-size.cjs: single source of truth for LF-normalized byte
  counting (#683) + workflow enumeration, shared by the guard and the generator
  so they can never measure differently. Lives in scripts/ root (NOT scripts/lib/)
  because it is dev/CI-only tooling — scripts/lib/ is bundled into the installed
  runtime, scripts/ root is not, so this keeps it out of the shipped payload.
- scripts/update-size-baseline.cjs + npm run size:baseline: regenerate the
  snapshot (sorted keys, trailing newline, idempotent).
- tests/workflow-size-baseline.json: generated snapshot (88 workflows).
- tests/workflow-size-budget.test.cjs: import the shared counter (drops the
  duplicated local byteCount) and add the per-file baseline describe block.
- Tests for the helper, the shared module, and the generator (incl. round-trip
  and fault-injection cases).

Refs #1074. Part 1 of 3; PR 2 swaps enforcement, PR 3 covers the agent test.

* test(#1074): regenerate workflow baseline after Update-branch merge with next

The 'Update branch' merge (652a916b) pulled in next's update.md change (#1090)
without regenerating the snapshot, leaving the per-file baseline stale by one
file. Re-ran `npm run size:baseline` so the committed baseline matches the
merged workflow files.

Refs #1074.

---------

Co-authored-by: Tom Boucher <trekkie@nomorestars.com>
This commit is contained in:
Rezolv
2026-06-11 23:59:56 -04:00
committed by GitHub
parent b4a7eabaae
commit 74d7bc8239
9 changed files with 723 additions and 14 deletions

View File

@@ -103,6 +103,7 @@
"lint:docs": "node scripts/lint-docs-required.cjs",
"lint:legacy-name": "node scripts/lint-legacy-dir-name.cjs",
"ci:test-scope": "node scripts/ci-test-scope.cjs",
"size:baseline": "node scripts/update-size-baseline.cjs",
"changeset": "node scripts/changeset/new.cjs",
"changelog:render": "node scripts/changeset/cli.cjs render",
"test": "node scripts/run-tests.cjs",

View File

@@ -133,4 +133,104 @@ function assertTightCeiling({ label, actualMax, ceiling, grace, fail }) {
return { ok, slack };
}
module.exports = { assertWithinAllowlist, assertTightCeiling };
/**
* Assert that each artifact's measured size matches a committed per-file
* baseline snapshot. Growth, shrinkage, additions, and removals are each
* surfaced by name — there is no aggregate "max" that can mask one file's
* growth behind another file's size.
*
* Fails when, for the union of `current` and `baseline` keys:
* - `current[name] > baseline[name]` → GROWTH: the file grew past its recorded
* size. Regenerate the baseline and justify the growth in the PR (or extract
* the content lazily). This is the headline guard.
* - `current[name] < baseline[name]` → STALE: the file shrank but the baseline
* still records the old (larger) size. Regenerate to auto-tighten — the
* per-file analogue of `assertWithinAllowlist`'s stale-entry rule, so the
* snapshot can only ratchet downward.
* - name in `current` but not `baseline` → ADDED: a new artifact with no
* recorded baseline. Regenerate to record it.
* - name in `baseline` but not `current` → REMOVED: an orphaned baseline entry
* whose artifact no longer exists. Regenerate to drop it.
*
* ## Why per-file, not a tier max (issue #1074)
*
* A `max(group) within grace` ceiling only binds the single largest file in the
* group; every other file inherits that ceiling and can grow silently beneath
* it. Recording each file's exact size removes the masking blind spot — the
* same reason `assertWithinAllowlist` enforces on identity rather than a count
* (issue #597).
*
* @param {object} opts
* @param {string} opts.label - Human-readable guard name (used in messages).
* @param {Object<string, number>} opts.current - Measured sizes by name.
* @param {Object<string, number>} opts.baseline - Committed sizes by name.
* @param {function(string): void} opts.fail - Callback invoked once per
* non-empty violation category with a
* descriptive message. Injected so callers
* control the failure mode (assert.fail, a
* thrower, or a collector in unit tests).
* @param {string} [opts.updateHint] - Optional remediation hint appended to
* every failure message (e.g. the regen
* command).
* @returns {{ grown: Array<{name:string,from:number,to:number,delta:number}>,
* shrunk: Array<{name:string,from:number,to:number,delta:number}>,
* added: string[], removed: string[] }}
* Sorted-by-name breakdown of every difference.
*/
function assertFileBaseline({ label, current, baseline, fail, updateHint }) {
const currentNames = new Set(Object.keys(current));
const baselineNames = new Set(Object.keys(baseline));
const added = [...currentNames].filter((n) => !baselineNames.has(n)).sort();
const removed = [...baselineNames].filter((n) => !currentNames.has(n)).sort();
const grown = [];
const shrunk = [];
const shared = [...currentNames].filter((n) => baselineNames.has(n)).sort();
for (const name of shared) {
const from = baseline[name];
const to = current[name];
if (to > from) grown.push({ name, from, to, delta: to - from });
else if (to < from) shrunk.push({ name, from, to, delta: from - to });
}
const hint = updateHint ? `\n${updateHint}` : '';
if (grown.length > 0) {
const list = grown
.map((g) => ` - ${g.name}: ${g.from} → ${g.to} (+${g.delta})`)
.join('\n');
fail(
`[${label}] ${grown.length} file(s) grew past the committed baseline. ` +
`Regenerate the baseline and justify the growth in your PR, or extract the content lazily.\n${list}${hint}`
);
}
if (shrunk.length > 0) {
const list = shrunk
.map((s) => ` - ${s.name}: ${s.from} → ${s.to} (-${s.delta})`)
.join('\n');
fail(
`[${label}] ${shrunk.length} file(s) are SMALLER than the baseline — the snapshot is stale ` +
`and MUST be regenerated so the budget ratchets downward.\n${list}${hint}`
);
}
if (added.length > 0) {
const list = added.map((n) => ` - ${n}`).join('\n');
fail(
`[${label}] ${added.length} file(s) are not in the baseline — regenerate to record them.\n${list}${hint}`
);
}
if (removed.length > 0) {
const list = removed.map((n) => ` - ${n}`).join('\n');
fail(
`[${label}] ${removed.length} baseline entry(ies) no longer exist — regenerate to drop them.\n${list}${hint}`
);
}
return { grown, shrunk, added, removed };
}
module.exports = { assertWithinAllowlist, assertTightCeiling, assertFileBaseline };

View File

@@ -0,0 +1,58 @@
#!/usr/bin/env node
'use strict';
/**
* @file update-size-baseline.cjs
*
* Regenerates the committed per-file workflow size baseline
* (`tests/workflow-size-baseline.json`) from the current workflow files.
*
* Run via `npm run size:baseline` whenever a workflow file legitimately grows
* or shrinks. Growth must still be justified in the PR; this script only
* records the new reality so the CI guard (issue #1074) can diff against it.
*
* Idempotent: running it twice with no file changes produces no diff.
*/
const fs = require('fs');
const path = require('path');
const { measureWorkflows, WORKFLOWS_DIR } = require('./workflow-size.cjs');
const BASELINE_PATH = path.join(__dirname, '..', 'tests', 'workflow-size-baseline.json');
/**
* Serialize a size map to the on-disk baseline format: keys sorted, 2-space
* indent, trailing newline (so the file is a stable, minimal-diff artifact).
*
* @param {Object<string, number>} sizes
* @returns {string}
*/
function serializeBaseline(sizes) {
const sorted = {};
for (const key of Object.keys(sizes).sort()) sorted[key] = sizes[key];
return JSON.stringify(sorted, null, 2) + '\n';
}
/**
* Write the baseline file from the measured workflow sizes.
*
* @param {object} [opts]
* @param {string} [opts.dir] - Workflows dir to measure (default canonical).
* @param {string} [opts.outPath] - Baseline file to write (default canonical).
* @returns {{ outPath: string, count: number, content: string }}
*/
function generateBaseline({ dir = WORKFLOWS_DIR, outPath = BASELINE_PATH } = {}) {
const sizes = measureWorkflows(dir);
const content = serializeBaseline(sizes);
fs.writeFileSync(outPath, content);
return { outPath, count: Object.keys(sizes).length, content };
}
if (require.main === module) {
const { outPath, count } = generateBaseline();
process.stdout.write(
`Wrote ${count} workflow sizes to ${path.relative(process.cwd(), outPath)}\n`
);
}
module.exports = { generateBaseline, serializeBaseline, BASELINE_PATH };

68
scripts/workflow-size.cjs Normal file
View File

@@ -0,0 +1,68 @@
'use strict';
/**
* @file workflow-size.cjs
*
* Single source of truth for measuring workflow `.md` file sizes in bytes.
*
* Shared by `tests/workflow-size-budget.test.cjs` (the CI guard) and
* `scripts/update-size-baseline.cjs` (the baseline generator) so the two can
* never disagree on HOW a file is measured. A divergence between the generator
* and the guard would silently mis-record the baseline (issue #1074).
*/
const fs = require('fs');
const path = require('path');
const WORKFLOWS_DIR = path.join(__dirname, '..', 'gsd-core', 'workflows');
/**
* Byte size of a file, counted as on an LF (Unix) checkout.
*
* The size budget is calibrated against `wc -c` on a Unix (LF) checkout, but
* these `.md` files have no `eol=lf` in `.gitattributes`, so Windows checks
* them out as CRLF. Counting raw on-disk bytes there adds one byte per line,
* a Windows-only false positive that diverges from the LF calibration basis
* (issue #683). Stripping CR yields the same LF byte count on every platform.
* This is still a raw byte count (not a trailing-newline-stripping line count).
*
* @param {string} filePath - Absolute or relative path to the file.
* @returns {number} LF-normalized byte length.
*/
function lfByteCount(filePath) {
const content = fs.readFileSync(filePath, 'utf-8');
return Buffer.byteLength(content.replace(/\r\n/g, '\n'), 'utf-8');
}
/**
* List top-level workflow stems (filenames without the `.md` extension), sorted.
* Non-recursive by design: per-mode bodies under `workflows/<name>/modes/` and
* templates are NOT measured — only the always-loaded top-level workflows.
*
* @param {string} [dir] - Workflows directory (defaults to the canonical one).
* @returns {string[]} Sorted stems, e.g. `['autonomous', 'plan-phase', ...]`.
*/
function listWorkflowStems(dir = WORKFLOWS_DIR) {
return fs
.readdirSync(dir)
.filter((f) => f.endsWith('.md'))
.map((f) => f.replace(/\.md$/, ''))
.sort();
}
/**
* Measure every top-level workflow file, keyed by filename (`<stem>.md`).
*
* @param {string} [dir] - Workflows directory (defaults to the canonical one).
* @returns {Object<string, number>} Map of `<stem>.md` → LF byte size, with
* keys inserted in sorted order.
*/
function measureWorkflows(dir = WORKFLOWS_DIR) {
const out = {};
for (const stem of listWorkflowStems(dir)) {
out[`${stem}.md`] = lfByteCount(path.join(dir, `${stem}.md`));
}
return out;
}
module.exports = { WORKFLOWS_DIR, lfByteCount, listWorkflowStems, measureWorkflows };

View File

@@ -14,6 +14,7 @@ const assert = require('node:assert/strict');
const {
assertWithinAllowlist,
assertTightCeiling,
assertFileBaseline,
} = require('../scripts/lib/allowlist-ratchet.cjs');
// ─── Fake fail helper ────────────────────────────────────────────────────────
@@ -295,3 +296,150 @@ describe('assertTightCeiling', () => {
assert.ok(calls[0].includes('my-special-guard'), 'label should appear in message');
});
});
// ─── assertFileBaseline ──────────────────────────────────────────────────────
describe('assertFileBaseline', () => {
test('exact match: current equals baseline — fail never called', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'a.md': 100, 'b.md': 200 },
baseline: { 'a.md': 100, 'b.md': 200 },
fail,
});
assert.strictEqual(calls.length, 0, 'fail should not be called on an exact match');
assert.deepStrictEqual(result.grown, []);
assert.deepStrictEqual(result.shrunk, []);
assert.deepStrictEqual(result.added, []);
assert.deepStrictEqual(result.removed, []);
});
test('growth: a file larger than baseline — fail called, delta reported', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'a.md': 154, 'b.md': 200 },
baseline: { 'a.md': 100, 'b.md': 200 },
fail,
});
assert.strictEqual(calls.length, 1, 'fail should be called once for growth');
assert.ok(calls[0].includes('grew') || calls[0].includes('grow'), 'message should describe growth');
assert.ok(calls[0].includes('a.md'), 'message should name the grown file');
assert.ok(calls[0].includes('100') && calls[0].includes('154'), 'message should show from → to');
assert.ok(calls[0].includes('54'), 'message should show the +delta');
assert.deepStrictEqual(result.grown.map((g) => g.name), ['a.md']);
assert.strictEqual(result.grown[0].delta, 54);
});
test('shrink: a file smaller than baseline — fail called as stale (auto-tighten)', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'a.md': 80, 'b.md': 200 },
baseline: { 'a.md': 100, 'b.md': 200 },
fail,
});
assert.strictEqual(calls.length, 1, 'fail should be called once for a stale (shrunk) baseline');
assert.ok(/stale|smaller|shrank|shrunk/i.test(calls[0]), 'message should flag a stale/shrunk baseline');
assert.ok(calls[0].includes('a.md'), 'message should name the shrunk file');
assert.deepStrictEqual(result.shrunk.map((s) => s.name), ['a.md']);
assert.strictEqual(result.shrunk[0].delta, 20);
});
test('added: a file absent from baseline — fail called', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'a.md': 100, 'new.md': 50 },
baseline: { 'a.md': 100 },
fail,
});
assert.strictEqual(calls.length, 1, 'fail should be called once for an unbaselined new file');
assert.ok(/not in the baseline|new|missing/i.test(calls[0]), 'message should flag the unbaselined file');
assert.ok(calls[0].includes('new.md'), 'message should name the new file');
assert.deepStrictEqual(result.added, ['new.md']);
});
test('removed: a baseline entry with no current file — fail called', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'a.md': 100 },
baseline: { 'a.md': 100, 'gone.md': 70 },
fail,
});
assert.strictEqual(calls.length, 1, 'fail should be called once for an orphaned baseline entry');
assert.ok(/no longer exist|removed|orphan/i.test(calls[0]), 'message should flag the orphaned entry');
assert.ok(calls[0].includes('gone.md'), 'message should name the orphaned entry');
assert.deepStrictEqual(result.removed, ['gone.md']);
});
test('multiple categories at once — one fail per non-empty category', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'grow.md': 150, 'shrink.md': 50, 'new.md': 10 },
baseline: { 'grow.md': 100, 'shrink.md': 100, 'gone.md': 30 },
fail,
});
// grown(1) + shrunk(1) + added(1) + removed(1) = 4 categories
assert.strictEqual(calls.length, 4, 'one fail per non-empty category');
assert.deepStrictEqual(result.grown.map((g) => g.name), ['grow.md']);
assert.deepStrictEqual(result.shrunk.map((s) => s.name), ['shrink.md']);
assert.deepStrictEqual(result.added, ['new.md']);
assert.deepStrictEqual(result.removed, ['gone.md']);
});
test('updateHint appears in every failure message', () => {
const { fail, calls } = makeFail();
assertFileBaseline({
label: 'workflow-size',
current: { 'grow.md': 150, 'new.md': 10 },
baseline: { 'grow.md': 100, 'gone.md': 30 },
fail,
updateHint: 'Run `npm run size:baseline`',
});
assert.ok(calls.length >= 2, 'multiple categories should each fail');
for (const msg of calls) {
assert.ok(msg.includes('Run `npm run size:baseline`'), 'every message should carry the updateHint');
}
});
test('empty inputs — fail never called', () => {
const { fail, calls } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: {},
baseline: {},
fail,
});
assert.strictEqual(calls.length, 0);
assert.deepStrictEqual(result.grown, []);
assert.deepStrictEqual(result.shrunk, []);
assert.deepStrictEqual(result.added, []);
assert.deepStrictEqual(result.removed, []);
});
test('returned category lists are sorted by name', () => {
const { fail } = makeFail();
const result = assertFileBaseline({
label: 'workflow-size',
current: { 'z.md': 10, 'a.md': 10, 'm.md': 10 },
baseline: {},
fail,
});
assert.deepStrictEqual(result.added, ['a.md', 'm.md', 'z.md'], 'added should be sorted');
});
test('label appears in failure messages', () => {
const { fail, calls } = makeFail();
assertFileBaseline({
label: 'my-size-guard',
current: { 'a.md': 200 },
baseline: { 'a.md': 100 },
fail,
});
assert.ok(calls[0].includes('my-size-guard'), 'label should appear in message');
});
});

View File

@@ -0,0 +1,127 @@
'use strict';
/**
* Tests for scripts/update-size-baseline.cjs — the per-file workflow size
* baseline generator (issue #1074).
*/
const { test, describe, beforeEach, afterEach } = require('node:test');
const assert = require('node:assert/strict');
const fs = require('fs');
const os = require('node:os');
const path = require('path');
const {
generateBaseline,
serializeBaseline,
} = require('../scripts/update-size-baseline.cjs');
const { assertFileBaseline } = require('../scripts/lib/allowlist-ratchet.cjs');
const { measureWorkflows } = require('../scripts/workflow-size.cjs');
const { cleanup } = require('./helpers.cjs');
describe('serializeBaseline', () => {
test('keys are sorted and output ends with a trailing newline', () => {
const out = serializeBaseline({ 'z.md': 3, 'a.md': 1, 'm.md': 2 });
assert.ok(out.endsWith('\n'), 'must end with a trailing newline');
const keys = Object.keys(JSON.parse(out));
assert.deepStrictEqual(keys, ['a.md', 'm.md', 'z.md'], 'keys must be sorted');
});
test('is stable: same input serializes identically (minimal-diff artifact)', () => {
const input = { 'b.md': 2, 'a.md': 1 };
assert.strictEqual(serializeBaseline(input), serializeBaseline({ ...input }));
});
});
describe('generateBaseline', () => {
let dir;
let outPath;
beforeEach(() => {
dir = fs.mkdtempSync(path.join(os.tmpdir(), 'gsd-gen-baseline-'));
outPath = path.join(dir, 'baseline.json');
});
afterEach(() => cleanup(dir));
test('writes a baseline matching the measured workflow sizes', () => {
const wfDir = path.join(dir, 'workflows');
fs.mkdirSync(wfDir);
fs.writeFileSync(path.join(wfDir, 'one.md'), 'hello\n');
fs.writeFileSync(path.join(wfDir, 'two.md'), 'a longer body here\n');
const result = generateBaseline({ dir: wfDir, outPath });
assert.strictEqual(result.count, 2);
const written = JSON.parse(fs.readFileSync(outPath, 'utf-8'));
assert.deepStrictEqual(written, measureWorkflows(wfDir));
});
test('idempotent: a second run with no changes produces an identical file', () => {
const wfDir = path.join(dir, 'workflows');
fs.mkdirSync(wfDir);
fs.writeFileSync(path.join(wfDir, 'one.md'), 'hello\n');
generateBaseline({ dir: wfDir, outPath });
const first = fs.readFileSync(outPath, 'utf-8');
generateBaseline({ dir: wfDir, outPath });
const second = fs.readFileSync(outPath, 'utf-8');
assert.strictEqual(first, second, 'a no-op regeneration must not churn the file');
});
test('round-trip: a freshly generated baseline satisfies assertFileBaseline', () => {
const wfDir = path.join(dir, 'workflows');
fs.mkdirSync(wfDir);
fs.writeFileSync(path.join(wfDir, 'one.md'), 'hello\n');
fs.writeFileSync(path.join(wfDir, 'two.md'), 'world body\n');
generateBaseline({ dir: wfDir, outPath });
const baseline = JSON.parse(fs.readFileSync(outPath, 'utf-8'));
const calls = [];
assertFileBaseline({
label: 'roundtrip',
current: measureWorkflows(wfDir),
baseline,
fail: (m) => calls.push(m),
});
assert.deepStrictEqual(calls, [], 'a just-generated baseline must pass the guard with zero failures');
});
test('regeneration records growth after a workflow file grows', () => {
const wfDir = path.join(dir, 'workflows');
fs.mkdirSync(wfDir);
const wf = path.join(wfDir, 'one.md');
fs.writeFileSync(wf, 'small\n');
generateBaseline({ dir: wfDir, outPath });
// File grows; the OLD baseline should now flag growth...
fs.writeFileSync(wf, 'small\nplus several more bytes\n');
const oldBaseline = JSON.parse(fs.readFileSync(outPath, 'utf-8'));
const beforeRegen = [];
assertFileBaseline({
label: 'grow',
current: measureWorkflows(wfDir),
baseline: oldBaseline,
fail: (m) => beforeRegen.push(m),
});
assert.strictEqual(beforeRegen.length, 1, 'old baseline must flag the growth');
// ...and regenerating clears it.
generateBaseline({ dir: wfDir, outPath });
const newBaseline = JSON.parse(fs.readFileSync(outPath, 'utf-8'));
const afterRegen = [];
assertFileBaseline({
label: 'grow',
current: measureWorkflows(wfDir),
baseline: newBaseline,
fail: (m) => afterRegen.push(m),
});
assert.deepStrictEqual(afterRegen, [], 'regenerated baseline must pass');
});
test('throws when the workflows directory does not exist', () => {
assert.throws(
() => generateBaseline({ dir: path.join(dir, 'missing'), outPath }),
/ENOENT/
);
});
});

View File

@@ -0,0 +1,90 @@
{
"add-backlog.md": 7132,
"add-phase.md": 7203,
"add-tests.md": 16901,
"add-todo.md": 8952,
"ai-integration-phase.md": 14761,
"analyze-dependencies.md": 3887,
"audit-fix.md": 10988,
"audit-milestone.md": 17556,
"audit-uat.md": 7425,
"autonomous.md": 41838,
"check-todos.md": 9431,
"cleanup.md": 9897,
"code-review-fix.md": 23676,
"code-review.md": 31397,
"complete-milestone.md": 30228,
"debug.md": 13505,
"diagnose-issues.md": 12425,
"discovery-phase.md": 8651,
"discuss-phase-assumptions.md": 26984,
"discuss-phase-power.md": 11273,
"discuss-phase.md": 31423,
"do.md": 10068,
"docs-update.md": 54770,
"edit-phase.md": 12883,
"eval-review.md": 9923,
"execute-phase.md": 92885,
"execute-plan.md": 29980,
"explore.md": 10497,
"extract-learnings.md": 12849,
"fast.md": 4149,
"forensics.md": 12463,
"graduation.md": 11578,
"health.md": 11824,
"help.md": 1722,
"import.md": 14604,
"inbox.md": 14407,
"ingest-docs.md": 18336,
"insert-phase.md": 8943,
"list-phase-assumptions.md": 4305,
"list-workspaces.md": 5655,
"manager.md": 25937,
"map-codebase.md": 20360,
"milestone-summary.md": 11774,
"mvp-phase.md": 13582,
"new-milestone.md": 32422,
"new-project.md": 61690,
"new-workspace.md": 11254,
"next.md": 17868,
"node-repair.md": 4173,
"note.md": 6563,
"pause-work.md": 13654,
"plan-milestone-gaps.md": 11765,
"plan-phase.md": 93135,
"plan-review-convergence.md": 22949,
"plant-seed.md": 11741,
"pr-branch.md": 4994,
"profile-user.md": 20457,
"progress.md": 26647,
"quick.md": 46213,
"reapply-patches.md": 20393,
"remove-phase.md": 8469,
"remove-workspace.md": 7507,
"resume-project.md": 15288,
"review.md": 37079,
"scan.md": 7688,
"secure-phase.md": 12187,
"session-report.md": 4044,
"settings-advanced.md": 39330,
"settings-integrations.md": 15801,
"settings.md": 32133,
"ship.md": 20896,
"sketch-wrap-up.md": 14223,
"sketch.md": 19960,
"spec-phase.md": 15131,
"spike-wrap-up.md": 15092,
"spike.md": 24517,
"stats.md": 6718,
"sync-skills.md": 6125,
"thread.md": 12400,
"transition.md": 21787,
"ui-phase.md": 15477,
"ui-review.md": 11289,
"ultraplan-phase.md": 10468,
"undo.md": 10431,
"update.md": 21053,
"validate-phase.md": 10630,
"verify-phase.md": 28430,
"verify-work.md": 30614
}

View File

@@ -66,10 +66,12 @@ const assert = require('node:assert/strict');
const fs = require('fs');
const os = require('node:os');
const path = require('path');
const { assertTightCeiling } = require('../scripts/lib/allowlist-ratchet.cjs');
const { assertTightCeiling, assertFileBaseline } = require('../scripts/lib/allowlist-ratchet.cjs');
const { lfByteCount: byteCount, measureWorkflows } = require('../scripts/workflow-size.cjs');
const { cleanup } = require('./helpers.cjs');
const WORKFLOWS_DIR = path.join(__dirname, '..', 'gsd-core', 'workflows');
const BASELINE_PATH = path.join(__dirname, 'workflow-size-baseline.json');
// Grace band: maximum allowed slack (ceiling − actualMax) in BYTES before a
// ceiling is considered too loose. 3000 bytes ≈ the prior 60-line grace
@@ -127,18 +129,10 @@ function budgetFor(workflow) {
return { tier: 'DEFAULT', limit: DEFAULT_BUDGET };
}
function byteCount(filePath) {
// Count bytes as on an LF checkout, so the budget is platform-independent.
// The tier ceilings are calibrated against `wc -c` on a Unix (LF) checkout,
// but these .md files have no `eol=lf` in .gitattributes, so Windows checks
// them out as CRLF. Counting raw on-disk bytes there adds one byte per line,
// which fails CI on the high-water-mark file (execute-phase.md) on Windows
// ONLY — a false positive that diverges from the LF calibration basis (#683).
// Stripping CR yields the same LF byte count on every platform. Still a raw
// byte count (not the old trailing-newline-stripping lineCount()).
const content = fs.readFileSync(filePath, 'utf-8');
return Buffer.byteLength(content.replace(/\r\n/g, '\n'), 'utf-8');
}
// byteCount (LF-normalized, #683) is imported as `lfByteCount` from
// scripts/workflow-size.cjs — the single source of truth shared with the
// baseline generator so the guard and the snapshot can never measure
// differently. See the #683 regression test at the bottom of this file.
describe('SIZE: workflow byte-size budget', () => {
for (const workflow of ALL_WORKFLOWS) {
@@ -159,6 +153,31 @@ describe('SIZE: workflow byte-size budget', () => {
}
});
describe('SIZE: per-file workflow baseline (issue #1074)', () => {
// Per-file exact-size ratchet. Unlike the tier anti-creep block below — which
// only binds the single largest file in each tier — this guards EVERY
// workflow file by name against a committed snapshot
// (tests/workflow-size-baseline.json). Growth fails with the file and delta;
// shrinkage fails as a stale snapshot (regenerate to ratchet down). The fix
// for any failure is `npm run size:baseline` plus a PR justification for
// genuine growth (or lazy extraction). Runs side-by-side with the tier tests
// during the #1074 migration; the tier anti-creep block is removed in a
// follow-up once this is established.
test('every workflow file matches its committed baseline', () => {
const baseline = JSON.parse(fs.readFileSync(BASELINE_PATH, 'utf-8'));
const current = measureWorkflows();
assertFileBaseline({
label: 'workflow-size',
current,
baseline,
fail: assert.fail,
updateHint:
'Run `npm run size:baseline` to update tests/workflow-size-baseline.json, ' +
'then justify any growth in your PR (or extract content lazily — see workflows/discuss-phase/).',
});
});
});
describe('SIZE: tier anti-creep (tighten-only ceilings, issue #597)', () => {
// For each tier, compute the high-water mark (in bytes) across all files in
// that tier and assert the ceiling stays tight. Prevents budgets from

View File

@@ -0,0 +1,98 @@
'use strict';
/**
* Tests for scripts/workflow-size.cjs — the shared LF byte counter and
* workflow enumeration used by both the size guard and the baseline generator.
*/
const { test, describe, beforeEach, afterEach } = require('node:test');
const assert = require('node:assert/strict');
const fs = require('fs');
const os = require('node:os');
const path = require('path');
const {
lfByteCount,
listWorkflowStems,
measureWorkflows,
WORKFLOWS_DIR,
} = require('../scripts/workflow-size.cjs');
const { cleanup } = require('./helpers.cjs');
describe('lfByteCount', () => {
let dir;
beforeEach(() => {
dir = fs.mkdtempSync(path.join(os.tmpdir(), 'gsd-wfsize-'));
});
afterEach(() => cleanup(dir));
test('counts raw UTF-8 bytes of an LF file', () => {
const body = 'line one\nline two\n';
const p = path.join(dir, 'a.md');
fs.writeFileSync(p, body);
assert.strictEqual(lfByteCount(p), Buffer.byteLength(body, 'utf-8'));
});
test('CRLF and LF of the same logical content count identically (#683)', () => {
const body = 'alpha\nbeta\ngamma — multibyte dash\n';
const lf = path.join(dir, 'lf.md');
const crlf = path.join(dir, 'crlf.md');
fs.writeFileSync(lf, body);
fs.writeFileSync(crlf, body.replace(/\n/g, '\r\n'));
assert.strictEqual(lfByteCount(crlf), lfByteCount(lf));
});
test('multibyte characters count as their UTF-8 byte length', () => {
const body = '— 漢字 🚀\n'; // em-dash (3) + CJK (3 each) + emoji (4)
const p = path.join(dir, 'm.md');
fs.writeFileSync(p, body);
assert.strictEqual(lfByteCount(p), Buffer.byteLength(body, 'utf-8'));
});
test('empty file counts as zero bytes', () => {
const p = path.join(dir, 'empty.md');
fs.writeFileSync(p, '');
assert.strictEqual(lfByteCount(p), 0);
});
test('throws on a missing file (no silent zero)', () => {
assert.throws(() => lfByteCount(path.join(dir, 'nope.md')), /ENOENT/);
});
});
describe('listWorkflowStems / measureWorkflows', () => {
let dir;
beforeEach(() => {
dir = fs.mkdtempSync(path.join(os.tmpdir(), 'gsd-wfmeasure-'));
});
afterEach(() => cleanup(dir));
test('lists only .md stems, sorted, without extension', () => {
fs.writeFileSync(path.join(dir, 'zeta.md'), 'z');
fs.writeFileSync(path.join(dir, 'alpha.md'), 'a');
fs.writeFileSync(path.join(dir, 'notes.txt'), 'ignored');
assert.deepStrictEqual(listWorkflowStems(dir), ['alpha', 'zeta']);
});
test('does not recurse into subdirectories (modes/templates excluded)', () => {
fs.writeFileSync(path.join(dir, 'top.md'), 'x');
fs.mkdirSync(path.join(dir, 'modes'));
fs.writeFileSync(path.join(dir, 'modes', 'sub.md'), 'should not be counted');
assert.deepStrictEqual(listWorkflowStems(dir), ['top']);
assert.deepStrictEqual(Object.keys(measureWorkflows(dir)), ['top.md']);
});
test('measureWorkflows keys by <stem>.md with LF byte sizes', () => {
const body = 'hello\nworld\n';
fs.writeFileSync(path.join(dir, 'one.md'), body);
const sizes = measureWorkflows(dir);
assert.deepStrictEqual(sizes, { 'one.md': Buffer.byteLength(body, 'utf-8') });
});
test('canonical WORKFLOWS_DIR resolves to a real directory with workflows', () => {
assert.ok(fs.existsSync(WORKFLOWS_DIR), 'canonical workflows dir should exist');
const sizes = measureWorkflows();
assert.ok(Object.keys(sizes).length > 0, 'should measure at least one workflow');
assert.ok('plan-phase.md' in sizes, 'plan-phase.md should be among the measured workflows');
});
});