ludic/tools/test-grammar.js
Orkuncakilkaya 985f9ad8f2 Baseline: Ludic compiler + toolchain, Phase 1 syntax fixes complete
Self-hosted compiler (selfhost/*.ludic), runtime, examples, editor tooling,
and docs. Phase 1 of the syntax-redesign cohesion pass has landed:
edge-system fix, signature-query, when-alias, and the documentation truth-pass.
Suite green (14/14), C-free bootstrap fixpoint holds.

Co-Authored-By: Claude Opus 4.8 <noreply@anthropic.com>
2026-08-27 15:15:35 +03:00

168 lines
7.5 KiB
JavaScript

/* ============================================================================
* test-grammar.js — tokenise Ludic with the real TextMate engine.
*
* A TextMate grammar is a pile of Oniguruma regexes in a JSON file: nothing
* validates it, and a broken pattern shows up as "that word stopped being
* coloured" months later. This runs the grammar through vscode-textmate, the
* same engine VS Code uses, and asserts that specific spans get specific
* scopes — including a ```ludic fence inside Markdown, which is the part most
* likely to break silently.
*
* node tools/test-grammar.js
*
* Needs vscode-textmate and vscode-oniguruma; skips cleanly when they are not
* installed, so it never blocks the suite on a machine without npm.
* ==========================================================================*/
const fs = require('fs');
const path = require('path');
const ROOT = path.dirname(__dirname);
const SHARED = path.join(ROOT, 'tools', 'editors', 'shared');
let vsctm, oniguruma;
try {
const local = path.join(ROOT, 'tools', 'editors', 'vscode', 'node_modules');
const paths = [local, '/tmp/tmgrammar/node_modules', ...module.paths];
vsctm = require(require.resolve('vscode-textmate', { paths }));
oniguruma = require(require.resolve('vscode-oniguruma', { paths }));
} catch (e) {
console.log(' skip TextMate grammar (npm i vscode-textmate vscode-oniguruma to run)');
process.exit(0);
}
let passed = 0;
let failed = 0;
function ok(name) { passed++; console.log(` ok ${name}`); }
function bad(name, detail) { failed++; console.log(` FAIL ${name}\n ${detail}`); }
async function main() {
const wasmPath = require.resolve('vscode-oniguruma/release/onig.wasm', {
paths: [path.join(ROOT, 'tools', 'editors', 'vscode', 'node_modules'), '/tmp/tmgrammar/node_modules']
});
await oniguruma.loadWASM(fs.readFileSync(wasmPath).buffer);
const onigLib = Promise.resolve({
createOnigScanner: (sources) => new oniguruma.OnigScanner(sources),
createOnigString: (s) => new oniguruma.OnigString(s)
});
const grammars = {
'source.ludic': path.join(SHARED, 'ludic.tmLanguage.json'),
'markdown.ludic.codeblock': path.join(SHARED, 'ludic.markdown-injection.json')
};
const registry = new vsctm.Registry({
onigLib,
loadGrammar: (scope) => {
const file = grammars[scope];
if (!file) return Promise.resolve(null);
return Promise.resolve(vsctm.parseRawGrammar(fs.readFileSync(file, 'utf8'), file));
},
getInjections: (scope) =>
scope === 'text.html.markdown' || scope.startsWith('text.html.markdown')
? ['markdown.ludic.codeblock']
: undefined
});
const grammar = await registry.loadGrammar('source.ludic');
if (!grammar) {
bad('grammar loads', 'registry returned null');
return report();
}
ok('grammar loads without a regex error');
// Tokenise a line and return the scopes covering the first occurrence of a
// substring, so assertions read in terms of the source rather than offsets.
function scopesFor(line, needle, state = vsctm.INITIAL) {
const at = line.indexOf(needle);
if (at < 0) throw new Error(`${needle} not in ${line}`);
const result = grammar.tokenizeLine(line, state);
for (const t of result.tokens) {
if (t.startIndex <= at && at < t.endIndex) return t.scopes;
}
return [];
}
function has(scopes, prefix) { return scopes.some((s) => s.startsWith(prefix)); }
function expect(name, line, needle, prefix) {
try {
const scopes = scopesFor(line, needle);
if (has(scopes, prefix)) ok(name);
else bad(name, `${needle!== undefined ? JSON.stringify(needle) : ''} got [${scopes.join(' ')}], wanted ${prefix}`);
} catch (e) {
bad(name, e.message);
}
}
expect('comment', '# a note', '# a note', 'comment.line');
expect('string', 'const S: str = "hi"', '"hi"', 'string.quoted.double');
expect('hex literal', 'clear(0x0e0e16)', '0x0e0e16', 'constant.numeric.hex');
expect('fixed literal', 'let f = 1.5', '1.5', 'constant.numeric.fixed');
expect('declaration keyword', 'component Pos { x: int = 0 }', 'component', 'storage.type.component');
expect('component name', 'component Pos { x: int = 0 }', 'Pos', 'entity.name.type.component');
expect('primitive type', 'component Pos { x: int = 0 }', 'int', 'support.type.primitive');
expect('system name', 'system Move phase Update {', 'Move', 'entity.name.function.system');
expect('phase name', 'system Move phase Update {', 'Update', 'constant.language.phase');
expect('clause keyword', ' reads [Vel]', 'reads', 'keyword.other.clause');
expect('annotation', 'system Move @deterministic', '@deterministic', 'entity.name.function.decorator');
expect('builtin call', ' clear(0x000000)', 'clear', 'support.function.builtin');
expect('intrinsic call', ' mem_alloc(64)', 'mem_alloc', 'support.function.intrinsic');
expect('control keyword', ' if a == 1 { }', 'if', 'keyword.control');
expect('member access', ' p.x = 1', 'x', 'variable.other.property');
expect('query tag filter', 'for (p) in query [Pos, {Player}] {', 'Player', 'entity.name.type.tag');
expect('query component', 'for (p) in query [Pos, {Player}] {', 'Pos', 'entity.name.type.component');
expect('widget type', ' button id=Go text="go"', 'button', 'support.class.widget');
expect('widget prop', ' button id=Go text="go"', 'id', 'variable.parameter.widget');
expect('user function call', ' my_helper(1)', 'my_helper', 'entity.name.function');
expect('archetype in spawn', ' spawn Hero {', 'Hero', 'entity.name.type.archetype');
expect('scene in enter', ' enter Battle', 'Battle', 'entity.name.type.scene');
// A whole file must tokenise without the engine bailing out.
const sample = fs.readFileSync(path.join(ROOT, 'examples', 'menu.ludic'), 'utf8');
let state = vsctm.INITIAL;
let tokenCount = 0;
for (const line of sample.split('\n')) {
const r = grammar.tokenizeLine(line, state);
state = r.ruleStack;
tokenCount += r.tokens.length;
}
if (tokenCount > 200) ok(`tokenises examples/menu.ludic (${tokenCount} tokens)`);
else bad('tokenises examples/menu.ludic', `only ${tokenCount} tokens`);
// ---- the Markdown injection ---------------------------------------------
const mdGrammar = await registry.loadGrammar('markdown.ludic.codeblock');
if (!mdGrammar) {
bad('markdown injection grammar loads', 'registry returned null');
return report();
}
ok('markdown injection grammar loads without a regex error');
// Drive the injection directly: feed it the fence, then the body, and check
// the body is handed to source.ludic.
const lines = ['```ludic', 'component Pos { x: int = 0 }', '```'];
let mdState = vsctm.INITIAL;
const bodyScopes = [];
lines.forEach((line, i) => {
const r = mdGrammar.tokenizeLine(line, mdState);
mdState = r.ruleStack;
if (i === 1) for (const t of r.tokens) bodyScopes.push(t.scopes.join(' '));
});
const joined = bodyScopes.join(' | ');
if (joined.includes('meta.embedded.block.ludic')) ok('markdown fence body is embedded as Ludic');
else bad('markdown fence body is embedded as Ludic', joined || '(no tokens)');
if (joined.includes('storage.type.component')) ok('markdown fence body is highlighted by source.ludic');
else bad('markdown fence body is highlighted by source.ludic', joined || '(no tokens)');
report();
}
function report() {
console.log(` ${passed} passed, ${failed} failed`);
process.exit(failed ? 1 : 0);
}
main().catch((e) => {
console.log(` FAIL grammar test crashed\n ${e.stack}`);
process.exit(1);
});