9403edfa83
This implements support for named captures in
RegExp.prototype[@@replace] for when the replaceValue is not callable.
Named captures can be referenced from replacement strings by using the
"$<name>" syntax. A couple of examples:
let re = /(?<fst>.)(?<snd>.)/u;
"abcd".replace(re, "$<snd>$<fst>") // "bacd"
"abcd".replace(re, "$2$1") // "bacd" (numbered refs work as always)
"abcd".replace(re, "$<snd") // SyntaxError (unterminated named ref)
"abcd".replace(re, "$<42$1>") // "cd" (invalid name)
"abcd".replace(re, "$<thd>") // "cd" (non-existent name)
"abcd".replace(/(?<fst>.)|(?<snd>.)/u, "$<snd>") // "cd" (non-matched capture)
Support is currently behind the --harmony-regexp-named-captures flag.
BUG=v8:5437
Review-Url: https://codereview.chromium.org/2775303002
Cr-Original-Commit-Position: refs/heads/master@{#44171}
Committed: 17f13863b6
Review-Url: https://codereview.chromium.org/2775303002
Cr-Commit-Position: refs/heads/master@{#44182}
267 lines
9.8 KiB
JavaScript
267 lines
9.8 KiB
JavaScript
// Copyright 2015 the V8 project authors. All rights reserved.
|
|
// Use of this source code is governed by a BSD-style license that can be
|
|
// found in the LICENSE file.
|
|
|
|
// Flags: --harmony-regexp-named-captures
|
|
|
|
// Malformed named captures.
|
|
assertThrows("/(?<>a)/u"); // Empty name.
|
|
assertThrows("/(?<aa)/u"); // Unterminated name.
|
|
assertThrows("/(?<42a>a)/u"); // Name starting with digits.
|
|
assertThrows("/(?<:a>a)/u"); // Name starting with invalid char.
|
|
assertThrows("/(?<a:>a)/u"); // Name containing with invalid char.
|
|
assertThrows("/(?<a>a)(?<a>a)/u"); // Duplicate name.
|
|
assertThrows("/(?<a>a)(?<b>b)(?<a>a)/u"); // Duplicate name.
|
|
assertThrows("/\\k<a>/u"); // Invalid reference.
|
|
assertThrows("/(?<a>a)\\k<ab>/u"); // Invalid reference.
|
|
assertThrows("/(?<ab>a)\\k<a>/u"); // Invalid reference.
|
|
assertThrows("/\\k<a>(?<ab>a)/u"); // Invalid reference.
|
|
|
|
// Fallback behavior in non-unicode mode.
|
|
assertThrows("/(?<>a)/", SyntaxError);
|
|
assertThrows("/(?<aa)/", SyntaxError);
|
|
assertThrows("/(?<42a>a)/", SyntaxError);
|
|
assertThrows("/(?<:a>a)/", SyntaxError);
|
|
assertThrows("/(?<a:>a)/", SyntaxError);
|
|
assertThrows("/(?<a>a)(?<a>a)/", SyntaxError);
|
|
assertThrows("/(?<a>a)(?<b>b)(?<a>a)/", SyntaxError);
|
|
assertThrows("/(?<a>a)\\k<ab>/", SyntaxError);
|
|
assertThrows("/(?<ab>a)\\k<a>/", SyntaxError);
|
|
|
|
assertEquals(["k<a>"], "xxxk<a>xxx".match(/\k<a>/));
|
|
assertEquals(["k<a"], "xxxk<a>xxx".match(/\k<a/));
|
|
|
|
// Basic named groups.
|
|
assertEquals(["a", "a"], "bab".match(/(?<a>a)/u));
|
|
assertEquals(["a", "a"], "bab".match(/(?<a42>a)/u));
|
|
assertEquals(["a", "a"], "bab".match(/(?<_>a)/u));
|
|
assertEquals(["a", "a"], "bab".match(/(?<$>a)/u));
|
|
assertEquals(["bab", "a"], "bab".match(/.(?<$>a)./u));
|
|
assertEquals(["bab", "a", "b"], "bab".match(/.(?<a>a)(.)/u));
|
|
assertEquals(["bab", "a", "b"], "bab".match(/.(?<a>a)(?<b>.)/u));
|
|
assertEquals(["bab", "ab"], "bab".match(/.(?<a>\w\w)/u));
|
|
assertEquals(["bab", "bab"], "bab".match(/(?<a>\w\w\w)/u));
|
|
assertEquals(["bab", "ba", "b"], "bab".match(/(?<a>\w\w)(?<b>\w)/u));
|
|
|
|
assertEquals("bab".match(/(a)/u), "bab".match(/(?<a>a)/u));
|
|
assertEquals("bab".match(/(a)/u), "bab".match(/(?<a42>a)/u));
|
|
assertEquals("bab".match(/(a)/u), "bab".match(/(?<_>a)/u));
|
|
assertEquals("bab".match(/(a)/u), "bab".match(/(?<$>a)/u));
|
|
assertEquals("bab".match(/.(a)./u), "bab".match(/.(?<$>a)./u));
|
|
assertEquals("bab".match(/.(a)(.)/u), "bab".match(/.(?<a>a)(.)/u));
|
|
assertEquals("bab".match(/.(a)(.)/u), "bab".match(/.(?<a>a)(?<b>.)/u));
|
|
assertEquals("bab".match(/.(\w\w)/u), "bab".match(/.(?<a>\w\w)/u));
|
|
assertEquals("bab".match(/(\w\w\w)/u), "bab".match(/(?<a>\w\w\w)/u));
|
|
assertEquals("bab".match(/(\w\w)(\w)/u), "bab".match(/(?<a>\w\w)(?<b>\w)/u));
|
|
|
|
assertEquals(["bab", "b"], "bab".match(/(?<b>b).\1/u));
|
|
assertEquals(["baba", "b", "a"], "baba".match(/(.)(?<a>a)\1\2/u));
|
|
assertEquals(["baba", "b", "a", "b", "a"],
|
|
"baba".match(/(.)(?<a>a)(?<b>\1)(\2)/u));
|
|
assertEquals(["<a", "<"], "<a".match(/(?<lt><)a/u));
|
|
assertEquals([">a", ">"], ">a".match(/(?<gt>>)a/u));
|
|
|
|
// Named references.
|
|
assertEquals(["bab", "b"], "bab".match(/(?<b>.).\k<b>/u));
|
|
assertNull("baa".match(/(?<b>.).\k<b>/u));
|
|
|
|
// Nested groups.
|
|
assertEquals(["bab", "bab", "ab", "b"], "bab".match(/(?<a>.(?<b>.(?<c>.)))/u));
|
|
|
|
// Reference inside group.
|
|
assertEquals(["bab", "b"], "bab".match(/(?<a>\k<a>\w)../u));
|
|
|
|
// Reference before group.
|
|
assertEquals(["bab", "b"], "bab".match(/\k<a>(?<a>b)\w\k<a>/u));
|
|
assertEquals(["bab", "b", "a"], "bab".match(/(?<b>b)\k<a>(?<a>a)\k<b>/u));
|
|
|
|
// Reference properties.
|
|
assertEquals("a", /(?<a>a)(?<b>b)\k<a>/u.exec("aba").groups.a);
|
|
assertEquals("b", /(?<a>a)(?<b>b)\k<a>/u.exec("aba").groups.b);
|
|
assertEquals(undefined, /(?<a>a)(?<b>b)\k<a>/u.exec("aba").groups.c);
|
|
assertEquals(undefined, /(?<a>a)(?<b>b)\k<a>|(?<c>c)/u.exec("aba").groups.c);
|
|
|
|
// Unicode names.
|
|
assertEquals("a", /(?<π>a)/u.exec("bab").groups.π);
|
|
assertEquals("a", /(?<\u{03C0}>a)/u.exec("bab").groups.\u03C0);
|
|
assertEquals("a", /(?<$>a)/u.exec("bab").groups.$);
|
|
assertEquals("a", /(?<_>a)/u.exec("bab").groups._);
|
|
assertEquals("a", /(?<$𐒤>a)/u.exec("bab").groups.$𐒤);
|
|
assertEquals("a", /(?<_\u200C>a)/u.exec("bab").groups._\u200C);
|
|
assertEquals("a", /(?<_\u200D>a)/u.exec("bab").groups._\u200D);
|
|
assertEquals("a", /(?<ಠ_ಠ>a)/u.exec("bab").groups.ಠ_ಠ);
|
|
assertThrows('/(?<❤>a)/u', SyntaxError);
|
|
assertThrows('/(?<𐒤>a)/u', SyntaxError); // ID_Continue but not ID_Start.
|
|
|
|
// The '__proto__' property on the groups object.
|
|
assertEquals(undefined, /(?<a>.)/u.exec("a").groups.__proto__);
|
|
assertEquals("a", /(?<__proto__>a)/u.exec("a").groups.__proto__);
|
|
|
|
// @@replace with a callable replacement argument (no named captures).
|
|
{
|
|
let result = "abcd".replace(/(.)(.)/u, (match, fst, snd, offset, str) => {
|
|
assertEquals("ab", match);
|
|
assertEquals("a", fst);
|
|
assertEquals("b", snd);
|
|
assertEquals(0, offset);
|
|
assertEquals("abcd", str);
|
|
return `${snd}${fst}`;
|
|
});
|
|
assertEquals("bacd", result);
|
|
}
|
|
|
|
// @@replace with a callable replacement argument (global, named captures).
|
|
{
|
|
let i = 0;
|
|
let result = "abcd".replace(/(?<fst>.)(?<snd>.)/gu,
|
|
(match, fst, snd, offset, str, groups) => {
|
|
if (i == 0) {
|
|
assertEquals("ab", match);
|
|
assertEquals("a", groups.fst);
|
|
assertEquals("b", groups.snd);
|
|
assertEquals("a", fst);
|
|
assertEquals("b", snd);
|
|
assertEquals(0, offset);
|
|
assertEquals("abcd", str);
|
|
} else if (i == 1) {
|
|
assertEquals("cd", match);
|
|
assertEquals("c", groups.fst);
|
|
assertEquals("d", groups.snd);
|
|
assertEquals("c", fst);
|
|
assertEquals("d", snd);
|
|
assertEquals(2, offset);
|
|
assertEquals("abcd", str);
|
|
} else {
|
|
assertUnreachable();
|
|
}
|
|
i++;
|
|
return `${groups.snd}${groups.fst}`;
|
|
});
|
|
assertEquals("badc", result);
|
|
}
|
|
|
|
// @@replace with a callable replacement argument (non-global, named captures).
|
|
{
|
|
let result = "abcd".replace(/(?<fst>.)(?<snd>.)/u,
|
|
(match, fst, snd, offset, str, groups) => {
|
|
assertEquals("ab", match);
|
|
assertEquals("a", groups.fst);
|
|
assertEquals("b", groups.snd);
|
|
assertEquals("a", fst);
|
|
assertEquals("b", snd);
|
|
assertEquals(0, offset);
|
|
assertEquals("abcd", str);
|
|
return `${groups.snd}${groups.fst}`;
|
|
});
|
|
assertEquals("bacd", result);
|
|
}
|
|
|
|
function toSlowMode(re) {
|
|
re.exec = (str) => RegExp.prototype.exec.call(re, str);
|
|
return re;
|
|
}
|
|
|
|
// @@replace with a callable replacement argument (slow, global,
|
|
// named captures).
|
|
{
|
|
let i = 0;
|
|
let re = toSlowMode(/(?<fst>.)(?<snd>.)/gu);
|
|
let result = "abcd".replace(re, (match, fst, snd, offset, str, groups) => {
|
|
if (i == 0) {
|
|
assertEquals("ab", match);
|
|
assertEquals("a", groups.fst);
|
|
assertEquals("b", groups.snd);
|
|
assertEquals("a", fst);
|
|
assertEquals("b", snd);
|
|
assertEquals(0, offset);
|
|
assertEquals("abcd", str);
|
|
} else if (i == 1) {
|
|
assertEquals("cd", match);
|
|
assertEquals("c", groups.fst);
|
|
assertEquals("d", groups.snd);
|
|
assertEquals("c", fst);
|
|
assertEquals("d", snd);
|
|
assertEquals(2, offset);
|
|
assertEquals("abcd", str);
|
|
} else {
|
|
assertUnreachable();
|
|
}
|
|
i++;
|
|
return `${groups.snd}${groups.fst}`;
|
|
});
|
|
assertEquals("badc", result);
|
|
}
|
|
|
|
// @@replace with a callable replacement argument (slow, non-global,
|
|
// named captures).
|
|
{
|
|
let re = toSlowMode(/(?<fst>.)(?<snd>.)/u);
|
|
let result = "abcd".replace(re, (match, fst, snd, offset, str, groups) => {
|
|
assertEquals("ab", match);
|
|
assertEquals("a", groups.fst);
|
|
assertEquals("b", groups.snd);
|
|
assertEquals("a", fst);
|
|
assertEquals("b", snd);
|
|
assertEquals(0, offset);
|
|
assertEquals("abcd", str);
|
|
return `${groups.snd}${groups.fst}`;
|
|
});
|
|
assertEquals("bacd", result);
|
|
}
|
|
|
|
// @@replace with a string replacement argument (no named captures).
|
|
{
|
|
let re = /(.)(.)/u;
|
|
assertEquals("$<snd>$<fst>cd", "abcd".replace(re, "$<snd>$<fst>"));
|
|
assertEquals("bacd", "abcd".replace(re, "$2$1"));
|
|
assertEquals("$<sndcd", "abcd".replace(re, "$<snd"));
|
|
assertEquals("$<42a>cd", "abcd".replace(re, "$<42$1>"));
|
|
assertEquals("$<thd>cd", "abcd".replace(re, "$<thd>"));
|
|
assertEquals("$<a>cd", "abcd".replace(re, "$<$1>"));
|
|
}
|
|
|
|
// @@replace with a string replacement argument (global, named captures).
|
|
{
|
|
let re = /(?<fst>.)(?<snd>.)/gu;
|
|
assertEquals("badc", "abcd".replace(re, "$<snd>$<fst>"));
|
|
assertEquals("badc", "abcd".replace(re, "$2$1"));
|
|
assertThrows(() => "abcd".replace(re, "$<snd"), SyntaxError);
|
|
assertEquals("", "abcd".replace(re, "$<42$1>"));
|
|
assertEquals("", "abcd".replace(re, "$<thd>"));
|
|
assertEquals("", "abcd".replace(re, "$<$1>"));
|
|
}
|
|
|
|
// @@replace with a string replacement argument (non-global, named captures).
|
|
{
|
|
let re = /(?<fst>.)(?<snd>.)/u;
|
|
assertEquals("bacd", "abcd".replace(re, "$<snd>$<fst>"));
|
|
assertEquals("bacd", "abcd".replace(re, "$2$1"));
|
|
assertThrows(() => "abcd".replace(re, "$<snd"), SyntaxError);
|
|
assertEquals("cd", "abcd".replace(re, "$<42$1>"));
|
|
assertEquals("cd", "abcd".replace(re, "$<thd>"));
|
|
assertEquals("cd", "abcd".replace(re, "$<$1>"));
|
|
}
|
|
|
|
// @@replace with a string replacement argument (slow, global, named captures).
|
|
{
|
|
let re = toSlowMode(/(?<fst>.)(?<snd>.)/gu);
|
|
assertEquals("badc", "abcd".replace(re, "$<snd>$<fst>"));
|
|
assertEquals("badc", "abcd".replace(re, "$2$1"));
|
|
assertThrows(() => "abcd".replace(re, "$<snd"), SyntaxError);
|
|
assertEquals("", "abcd".replace(re, "$<42$1>"));
|
|
assertEquals("", "abcd".replace(re, "$<thd>"));
|
|
assertEquals("", "abcd".replace(re, "$<$1>"));
|
|
}
|
|
|
|
// @@replace with a string replacement argument (slow, non-global,
|
|
// named captures).
|
|
{
|
|
let re = toSlowMode(/(?<fst>.)(?<snd>.)/u);
|
|
assertEquals("bacd", "abcd".replace(re, "$<snd>$<fst>"));
|
|
assertEquals("bacd", "abcd".replace(re, "$2$1"));
|
|
assertThrows(() => "abcd".replace(re, "$<snd"), SyntaxError);
|
|
assertEquals("cd", "abcd".replace(re, "$<42$1>"));
|
|
assertEquals("cd", "abcd".replace(re, "$<thd>"));
|
|
assertEquals("cd", "abcd".replace(re, "$<$1>"));
|
|
}
|