JavaCode/java-source/src/main/java/util/regex/Pattern.java at master · jxwsy/JavaCode

History

5856 lines (5557 loc) · 215 KB

Raw

100

101

102

103

104

105

106

107

108

109

110

111

112

113

114

115

116

117

118

119

120

121

122

123

124

125

126

127

128

129

130

131

132

133

134

135

136

137

138

139

140

141

142

143

144

145

146

147

148

149

150

151

152

153

154

155

156

157

158

159

160

161

162

163

164

165

166

167

168

169

170

171

172

173

174

175

176

177

178

179

180

181

182

183

184

185

186

187

188

189

190

191

192

193

194

195

196

197

198

199

200

201

202

203

204

205

206

207

208

209

210

211

212

213

214

215

216

217

218

219

220

221

222

223

224

225

226

227

228

229

230

231

232

233

234

235

236

237

238

239

240

241

242

243

244

245

246

247

248

249

250

251

252

253

254

255

256

257

258

259

260

261

262

263

264

265

266

267

268

269

270

271

272

273

274

275

276

277

278

279

280

281

282

283

284

285

286

287

288

289

290

291

292

293

294

295

296

297

298

299

300

301

302

303

304

305

306

307

308

309

310

311

312

313

314

315

316

317

318

319

320

321

322

323

324

325

326

327

328

329

330

331

332

333

334

335

336

337

338

339

340

341

342

343

344

345

346

347

348

349

350

351

352

353

354

355

356

357

358

359

360

361

362

363

364

365

366

367

368

369

370

371

372

373

374

375

376

377

378

379

380

381

382

383

384

385

386

387

388

389

390

391

392

393

394

395

396

397

398

399

400

401

402

403

404

405

406

407

408

409

410

411

412

413

414

415

416

417

418

419

420

421

422

423

424

425

426

427

428

429

430

431

432

433

434

435

436

437

438

439

440

441

442

443

444

445

446

447

448

449

450

451

452

453

454

455

456

457

458

459

460

461

462

463

464

465

466

467

468

469

470

471

472

473

474

475

476

477

478

479

480

481

482

483

484

485

486

487

488

489

490

491

492

493

494

495

496

497

498

499

500

501

502

503

504

505

506

507

508

509

510

511

512

513

514

515

516

517

518

519

520

521

522

523

524

525

526

527

528

529

530

531

532

533

534

535

536

537

538

539

540

541

542

543

544

545

546

547

548

549

550

551

552

553

554

555

556

557

558

559

560

561

562

563

564

565

566

567

568

569

570

571

572

573

574

575

576

577

578

579

580

581

582

583

584

585

586

587

588

589

590

591

592

593

594

595

596

597

598

599

600

601

602

603

604

605

606

607

608

609

610

611

612

613

614

615

616

617

618

619

620

621

622

623

624

625

626

627

628

629

630

631

632

633

634

635

636

637

638

639

640

641

642

643

644

645

646

647

648

649

650

651

652

653

654

655

656

657

658

659

660

661

662

663

664

665

666

667

668

669

670

671

672

673

674

675

676

677

678

679

680

681

682

683

684

685

686

687

688

689

690

691

692

693

694

695

696

697

698

699

700

701

702

703

704

705

706

707

708

709

710

711

712

713

714

715

716

717

718

719

720

721

722

723

724

725

726

727

728

729

730

731

732

733

734

735

736

737

738

739

740

741

742

743

744

745

746

747

748

749

750

751

752

753

754

755

756

757

758

759

760

761

762

763

764

765

766

767

768

769

770

771

772

773

774

775

776

777

778

779

780

781

782

783

784

785

786

787

788

789

790

791

792

793

794

795

796

797

798

799

800

801

802

803

804

805

806

807

808

809

810

811

812

813

814

815

816

817

818

819

820

821

822

823

824

825

826

827

828

829

830

831

832

833

834

835

836

837

838

839

840

841

842

843

844

845

846

847

848

849

850

851

852

853

854

855

856

857

858

859

860

861

862

863

864

865

866

867

868

869

870

871

872

873

874

875

876

877

878

879

880

881

882

883

884

885

886

887

888

889

890

891

892

893

894

895

896

897

898

899

900

901

902

903

904

905

906

907

908

909

910

911

912

913

914

915

916

917

918

919

920

921

922

923

924

925

926

927

928

929

930

931

932

933

934

935

936

937

938

939

940

941

942

943

944

945

946

947

948

949

950

951

952

953

954

955

956

957

958

959

960

961

962

963

964

965

966

967

968

969

970

971

972

973

974

975

976

977

978

979

980

981

982

983

984

985

986

987

988

989

990

991

992

993

994

995

996

997

998

999

1000

* ORACLE PROPRIETARY/CONFIDENTIAL. Use is subject to license terms.

package java.util.regex;

import java.text.Normalizer;

import java.util.Locale;

import java.util.Iterator;

import java.util.Map;

import java.util.ArrayList;

import java.util.HashMap;

import java.util.Arrays;

import java.util.NoSuchElementException;

import java.util.Spliterator;

import java.util.Spliterators;

import java.util.function.Predicate;

import java.util.stream.Stream;

import java.util.stream.StreamSupport;

/**

* A compiled representation of a regular expression.

* A regular expression, specified as a string, must first be compiled into

* an instance of this class. The resulting pattern can then be used to create

* a {@link Matcher} object that can match arbitrary {@linkplain

* java.lang.CharSequence character sequences} against the regular

* expression. All of the state involved in performing a match resides in the

* matcher, so many matchers can share the same pattern.

* A typical invocation sequence is thus

* <blockquote><pre>

* Pattern p = Pattern.{@link #compile compile}("a*b");

* Matcher m = p.{@link #matcher matcher}("aaaaab");

* boolean b = m.{@link Matcher#matches matches}();</pre></blockquote>

* A {@link #matches matches} method is defined by this class as a

* convenience for when a regular expression is used just once. This method

* compiles an expression and matches an input sequence against it in a single

* invocation. The statement

* <blockquote><pre>

* boolean b = Pattern.matches("a*b", "aaaaab");</pre></blockquote>

* is equivalent to the three statements above, though for repeated matches it

* is less efficient since it does not allow the compiled pattern to be reused.

* Instances of this class are immutable and are safe for use by multiple

* concurrent threads. Instances of the {@link Matcher} class are not safe for

* such use.

* <h3><a name="sum">Summary of regular-expression constructs</a></h3>

* <table border="0" cellpadding="1" cellspacing="0"

* summary="Regular expression constructs, and what they match">

* <tr align="left">

* <th align="left" id="construct">Construct</th>

* <th align="left" id="matches">Matches</th>

* </tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="characters">Characters</th></tr>

* <tr><td valign="top" headers="construct characters">x</td>

* <td headers="matches">The character x</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\\</tt></td>

* <td headers="matches">The backslash character</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\0</tt>n</td>

* <td headers="matches">The character with octal value <tt>0</tt>n

* (0 <tt><=</tt> n <tt><=</tt> 7)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\0</tt>nn</td>

* <td headers="matches">The character with octal value <tt>0</tt>nn

* (0 <tt><=</tt> n <tt><=</tt> 7)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\0</tt>mnn</td>

* <td headers="matches">The character with octal value <tt>0</tt>mnn

* (0 <tt><=</tt> m <tt><=</tt> 3,

* 0 <tt><=</tt> n <tt><=</tt> 7)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\x</tt>hh</td>

* <td headers="matches">The character with hexadecimal value <tt>0x</tt>hh</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\u</tt>hhhh</td>

* <td headers="matches">The character with hexadecimal value <tt>0x</tt>hhhh</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\x</tt>{h...h}</td>

* <td headers="matches">The character with hexadecimal value <tt>0x</tt>h...h

* ({@link java.lang.Character#MIN_CODE_POINT Character.MIN_CODE_POINT}

*  <= <tt>0x</tt>h...h <=

* {@link java.lang.Character#MAX_CODE_POINT Character.MAX_CODE_POINT})</td></tr>

* <tr><td valign="top" headers="matches"><tt>\t</tt></td>

* <td headers="matches">The tab character (<tt>'\u0009'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\n</tt></td>

* <td headers="matches">The newline (line feed) character (<tt>'\u000A'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\r</tt></td>

* <td headers="matches">The carriage-return character (<tt>'\u000D'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\f</tt></td>

* <td headers="matches">The form-feed character (<tt>'\u000C'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\a</tt></td>

* <td headers="matches">The alert (bell) character (<tt>'\u0007'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\e</tt></td>

* <td headers="matches">The escape character (<tt>'\u001B'</tt>)</td></tr>

* <tr><td valign="top" headers="construct characters"><tt>\c</tt>x</td>

* <td headers="matches">The control character corresponding to x</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="classes">Character classes</th></tr>

* <tr><td valign="top" headers="construct classes">{@code [abc]}</td>

* <td headers="matches">{@code a}, {@code b}, or {@code c} (simple class)</td></tr>

* <tr><td valign="top" headers="construct classes">{@code [^abc]}</td>

* <td headers="matches">Any character except {@code a}, {@code b}, or {@code c} (negation)</td></tr>

* <tr><td valign="top" headers="construct classes">{@code [a-zA-Z]}</td>

* <td headers="matches">{@code a} through {@code z}

* or {@code A} through {@code Z}, inclusive (range)</td></tr>

* <tr><td valign="top" headers="construct classes">{@code [a-d[m-p]]}</td>

* <td headers="matches">{@code a} through {@code d},

* or {@code m} through {@code p}: {@code [a-dm-p]} (union)</td></tr>

* <tr><td valign="top" headers="construct classes">{@code [a-z&&[def]]}</td>

* <td headers="matches">{@code d}, {@code e}, or {@code f} (intersection)</tr>

* <tr><td valign="top" headers="construct classes">{@code [a-z&&[^bc]]}</td>

* <td headers="matches">{@code a} through {@code z},

* except for {@code b} and {@code c}: {@code [ad-z]} (subtraction)</td></tr>

* <tr><td valign="top" headers="construct classes">{@code [a-z&&[^m-p]]}</td>

* <td headers="matches">{@code a} through {@code z},

* and not {@code m} through {@code p}: {@code [a-lq-z]}(subtraction)</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="predef">Predefined character classes</th></tr>

* <tr><td valign="top" headers="construct predef"><tt>.</tt></td>

* <td headers="matches">Any character (may or may not match <a href="#lt">line terminators</a>)</td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\d</tt></td>

* <td headers="matches">A digit: <tt>[0-9]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\D</tt></td>

* <td headers="matches">A non-digit: <tt>[^0-9]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\h</tt></td>

* <td headers="matches">A horizontal whitespace character:

* <tt>[ \t\xA0\u1680\u180e\u2000-\u200a\u202f\u205f\u3000]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\H</tt></td>

* <td headers="matches">A non-horizontal whitespace character: <tt>[^\h]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\s</tt></td>

* <td headers="matches">A whitespace character: <tt>[ \t\n\x0B\f\r]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\S</tt></td>

* <td headers="matches">A non-whitespace character: <tt>[^\s]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\v</tt></td>

* <td headers="matches">A vertical whitespace character: <tt>[\n\x0B\f\r\x85\u2028\u2029]</tt>

* </td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\V</tt></td>

* <td headers="matches">A non-vertical whitespace character: <tt>[^\v]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\w</tt></td>

* <td headers="matches">A word character: <tt>[a-zA-Z_0-9]</tt></td></tr>

* <tr><td valign="top" headers="construct predef"><tt>\W</tt></td>

* <td headers="matches">A non-word character: <tt>[^\w]</tt></td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="posix">POSIX character classes (US-ASCII only)</th></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Lower}}</td>

* <td headers="matches">A lower-case alphabetic character: {@code [a-z]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Upper}}</td>

* <td headers="matches">An upper-case alphabetic character:{@code [A-Z]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{ASCII}}</td>

* <td headers="matches">All ASCII:{@code [\x00-\x7F]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Alpha}}</td>

* <td headers="matches">An alphabetic character:{@code [\p{Lower}\p{Upper}]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Digit}}</td>

* <td headers="matches">A decimal digit: {@code [0-9]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Alnum}}</td>

* <td headers="matches">An alphanumeric character:{@code [\p{Alpha}\p{Digit}]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Punct}}</td>

* <td headers="matches">Punctuation: One of {@code !"#$%&'()*+,-./:;<=>?@[\]^_`{|}~}</td></tr>

* <!-- {@code [\!"#\$%&'\*\+,\-\./:;\<=\>\?@\[\\\]\^_`\{\|\}~]}

* {@code [\X21-\X2F\X31-\X40\X5B-\X60\X7B-\X7E]} -->

* <tr><td valign="top" headers="construct posix">{@code \p{Graph}}</td>

* <td headers="matches">A visible character: {@code [\p{Alnum}\p{Punct}]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Print}}</td>

* <td headers="matches">A printable character: {@code [\p{Graph}\x20]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Blank}}</td>

* <td headers="matches">A space or a tab: {@code [ \t]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Cntrl}}</td>

* <td headers="matches">A control character: {@code [\x00-\x1F\x7F]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{XDigit}}</td>

* <td headers="matches">A hexadecimal digit: {@code [0-9a-fA-F]}</td></tr>

* <tr><td valign="top" headers="construct posix">{@code \p{Space}}</td>

* <td headers="matches">A whitespace character: {@code [ \t\n\x0B\f\r]}</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2">java.lang.Character classes (simple <a href="#jcc">java character type</a>)</th></tr>

* <tr><td valign="top"><tt>\p{javaLowerCase}</tt></td>

* <td>Equivalent to java.lang.Character.isLowerCase()</td></tr>

* <tr><td valign="top"><tt>\p{javaUpperCase}</tt></td>

* <td>Equivalent to java.lang.Character.isUpperCase()</td></tr>

* <tr><td valign="top"><tt>\p{javaWhitespace}</tt></td>

* <td>Equivalent to java.lang.Character.isWhitespace()</td></tr>

* <tr><td valign="top"><tt>\p{javaMirrored}</tt></td>

* <td>Equivalent to java.lang.Character.isMirrored()</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="unicode">Classes for Unicode scripts, blocks, categories and binary properties</th></tr>

* <tr><td valign="top" headers="construct unicode">{@code \p{IsLatin}}</td>

* <td headers="matches">A Latin script character (<a href="#usc">script</a>)</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code \p{InGreek}}</td>

* <td headers="matches">A character in the Greek block (<a href="#ubc">block</a>)</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code \p{Lu}}</td>

* <td headers="matches">An uppercase letter (<a href="#ucc">category</a>)</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code \p{IsAlphabetic}}</td>

* <td headers="matches">An alphabetic character (<a href="#ubpc">binary property</a>)</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code \p{Sc}}</td>

* <td headers="matches">A currency symbol</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code \P{InGreek}}</td>

* <td headers="matches">Any character except one in the Greek block (negation)</td></tr>

* <tr><td valign="top" headers="construct unicode">{@code [\p{L}&&[^\p{Lu}]]}</td>

* <td headers="matches">Any letter except an uppercase letter (subtraction)</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="bounds">Boundary matchers</th></tr>

* <tr><td valign="top" headers="construct bounds"><tt>^</tt></td>

* <td headers="matches">The beginning of a line</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>$</tt></td>

* <td headers="matches">The end of a line</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\b</tt></td>

* <td headers="matches">A word boundary</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\B</tt></td>

* <td headers="matches">A non-word boundary</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\A</tt></td>

* <td headers="matches">The beginning of the input</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\G</tt></td>

* <td headers="matches">The end of the previous match</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\Z</tt></td>

* <td headers="matches">The end of the input but for the final

* <a href="#lt">terminator</a>, if any</td></tr>

* <tr><td valign="top" headers="construct bounds"><tt>\z</tt></td>

* <td headers="matches">The end of the input</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="lineending">Linebreak matcher</th></tr>

* <tr><td valign="top" headers="construct lineending"><tt>\R</tt></td>

* <td headers="matches">Any Unicode linebreak sequence, is equivalent to

* <tt>\u000D\u000A|[\u000A\u000B\u000C\u000D\u0085\u2028\u2029]

* </tt></td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="greedy">Greedy quantifiers</th></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>?</tt></td>

* <td headers="matches">X, once or not at all</td></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>*</tt></td>

* <td headers="matches">X, zero or more times</td></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>+</tt></td>

* <td headers="matches">X, one or more times</td></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>{</tt>n<tt>}</tt></td>

* <td headers="matches">X, exactly n times</td></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>{</tt>n<tt>,}</tt></td>

* <td headers="matches">X, at least n times</td></tr>

* <tr><td valign="top" headers="construct greedy">X<tt>{</tt>n<tt>,</tt>m<tt>}</tt></td>

* <td headers="matches">X, at least n but not more than m times</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="reluc">Reluctant quantifiers</th></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>??</tt></td>

* <td headers="matches">X, once or not at all</td></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>*?</tt></td>

* <td headers="matches">X, zero or more times</td></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>+?</tt></td>

* <td headers="matches">X, one or more times</td></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>{</tt>n<tt>}?</tt></td>

* <td headers="matches">X, exactly n times</td></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>{</tt>n<tt>,}?</tt></td>

* <td headers="matches">X, at least n times</td></tr>

* <tr><td valign="top" headers="construct reluc">X<tt>{</tt>n<tt>,</tt>m<tt>}?</tt></td>

* <td headers="matches">X, at least n but not more than m times</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="poss">Possessive quantifiers</th></tr>

* <tr><td valign="top" headers="construct poss">X<tt>?+</tt></td>

* <td headers="matches">X, once or not at all</td></tr>

* <tr><td valign="top" headers="construct poss">X<tt>*+</tt></td>

* <td headers="matches">X, zero or more times</td></tr>

* <tr><td valign="top" headers="construct poss">X<tt>++</tt></td>

* <td headers="matches">X, one or more times</td></tr>

* <tr><td valign="top" headers="construct poss">X<tt>{</tt>n<tt>}+</tt></td>

* <td headers="matches">X, exactly n times</td></tr>

* <tr><td valign="top" headers="construct poss">X<tt>{</tt>n<tt>,}+</tt></td>

* <td headers="matches">X, at least n times</td></tr>

* <tr><td valign="top" headers="construct poss">X<tt>{</tt>n<tt>,</tt>m<tt>}+</tt></td>

* <td headers="matches">X, at least n but not more than m times</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="logical">Logical operators</th></tr>

* <tr><td valign="top" headers="construct logical">XY</td>

* <td headers="matches">X followed by Y</td></tr>

* <tr><td valign="top" headers="construct logical">X<tt>|</tt>Y</td>

* <td headers="matches">Either X or Y</td></tr>

* <tr><td valign="top" headers="construct logical"><tt>(</tt>X<tt>)</tt></td>

* <td headers="matches">X, as a <a href="#cg">capturing group</a></td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="backref">Back references</th></tr>

* <tr><td valign="bottom" headers="construct backref"><tt>\</tt>n</td>

* <td valign="bottom" headers="matches">Whatever the nth

* <a href="#cg">capturing group</a> matched</td></tr>

* <tr><td valign="bottom" headers="construct backref"><tt>\</tt>k<name></td>

* <td valign="bottom" headers="matches">Whatever the

* <a href="#groupname">named-capturing group</a> "name" matched</td></tr>

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="quot">Quotation</th></tr>

* <tr><td valign="top" headers="construct quot"><tt>\</tt></td>

* <td headers="matches">Nothing, but quotes the following character</td></tr>

* <tr><td valign="top" headers="construct quot"><tt>\Q</tt></td>

* <td headers="matches">Nothing, but quotes all characters until <tt>\E</tt></td></tr>

* <tr><td valign="top" headers="construct quot"><tt>\E</tt></td>

* <td headers="matches">Nothing, but ends quoting started by <tt>\Q</tt></td></tr>

*

* <tr><th> </th></tr>

* <tr align="left"><th colspan="2" id="special">Special constructs (named-capturing and non-capturing)</th></tr>

* <tr><td valign="top" headers="construct special"><tt>(?<<a href="#groupname">name</a>></tt>X<tt>)</tt></td>

* <td headers="matches">X, as a named-capturing group</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?:</tt>X<tt>)</tt></td>

* <td headers="matches">X, as a non-capturing group</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?idmsuxU-idmsuxU) </tt></td>

* <td headers="matches">Nothing, but turns match flags <a href="#CASE_INSENSITIVE">i</a>

* <a href="#UNIX_LINES">d</a> <a href="#MULTILINE">m</a> <a href="#DOTALL">s</a>

* <a href="#UNICODE_CASE">u</a> <a href="#COMMENTS">x</a> <a href="#UNICODE_CHARACTER_CLASS">U</a>

* on - off</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?idmsux-idmsux:</tt>X<tt>)</tt>  </td>

* <td headers="matches">X, as a <a href="#cg">non-capturing group</a> with the

* given flags <a href="#CASE_INSENSITIVE">i</a> <a href="#UNIX_LINES">d</a>

* <a href="#MULTILINE">m</a> <a href="#DOTALL">s</a> <a href="#UNICODE_CASE">u</a >

* <a href="#COMMENTS">x</a> on - off</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?=</tt>X<tt>)</tt></td>

* <td headers="matches">X, via zero-width positive lookahead</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?!</tt>X<tt>)</tt></td>

* <td headers="matches">X, via zero-width negative lookahead</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?<=</tt>X<tt>)</tt></td>

* <td headers="matches">X, via zero-width positive lookbehind</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?<!</tt>X<tt>)</tt></td>

* <td headers="matches">X, via zero-width negative lookbehind</td></tr>

* <tr><td valign="top" headers="construct special"><tt>(?></tt>X<tt>)</tt></td>

* <td headers="matches">X, as an independent, non-capturing group</td></tr>

* </table>

* <hr>

* <h3><a name="bs">Backslashes, escapes, and quoting</a></h3>

* The backslash character (<tt>'\'</tt>) serves to introduce escaped

* constructs, as defined in the table above, as well as to quote characters

* that otherwise would be interpreted as unescaped constructs. Thus the

* expression <tt>\\</tt> matches a single backslash and <tt>\{</tt> matches a

* left brace.

* It is an error to use a backslash prior to any alphabetic character that

* does not denote an escaped construct; these are reserved for future

* extensions to the regular-expression language. A backslash may be used

* prior to a non-alphabetic character regardless of whether that character is

* part of an unescaped construct.

* Backslashes within string literals in Java source code are interpreted

* as required by

* <cite>The Java™ Language Specification</cite>

* as either Unicode escapes (section 3.3) or other character escapes (section 3.10.6)

* It is therefore necessary to double backslashes in string

* literals that represent regular expressions to protect them from

* interpretation by the Java bytecode compiler. The string literal

* <tt>"\b"</tt>, for example, matches a single backspace character when

* interpreted as a regular expression, while <tt>"\\b"</tt> matches a

* word boundary. The string literal <tt>"$hello$"</tt> is illegal

* and leads to a compile-time error; in order to match the string

* <tt>(hello)</tt> the string literal <tt>"\$hello\$"</tt>

* must be used.

* <h3><a name="cc">Character Classes</a></h3>

* Character classes may appear within other character classes, and

* may be composed by the union operator (implicit) and the intersection

* operator (<tt>&&</tt>).

* The union operator denotes a class that contains every character that is

* in at least one of its operand classes. The intersection operator

* denotes a class that contains every character that is in both of its

* operand classes.

* The precedence of character-class operators is as follows, from

* highest to lowest:

* <blockquote><table border="0" cellpadding="1" cellspacing="0"

* summary="Precedence of character class operators.">

* <tr><th>1    </th>

* <td>Literal escape    </td>

* <td><tt>\x</tt></td></tr>

* <tr><th>2    </th>

* <td>Grouping</td>

* <td><tt>[...]</tt></td></tr>

* <tr><th>3    </th>

* <td>Range</td>

* <td><tt>a-z</tt></td></tr>

* <tr><th>4    </th>

* <td>Union</td>

* <td><tt>[a-e][i-u]</tt></td></tr>

* <tr><th>5    </th>

* <td>Intersection</td>

* <td>{@code [a-z&&[aeiou]]}</td></tr>

* </table></blockquote>

* Note that a different set of metacharacters are in effect inside

* a character class than outside a character class. For instance, the

* regular expression <tt>.</tt> loses its special meaning inside a

* character class, while the expression <tt>-</tt> becomes a range

* forming metacharacter.

* <h3><a name="lt">Line terminators</a></h3>

* A line terminator is a one- or two-character sequence that marks

* the end of a line of the input character sequence. The following are

* recognized as line terminators:

* <ul>

* <li> A newline (line feed) character (<tt>'\n'</tt>),

* <li> A carriage-return character followed immediately by a newline

* character (<tt>"\r\n"</tt>),

* <li> A standalone carriage-return character (<tt>'\r'</tt>),

* <li> A next-line character (<tt>'\u0085'</tt>),

* <li> A line-separator character (<tt>'\u2028'</tt>), or

* <li> A paragraph-separator character (<tt>'\u2029</tt>).

* </ul>

* If {@link #UNIX_LINES} mode is activated, then the only line terminators

* recognized are newline characters.

* The regular expression <tt>.</tt> matches any character except a line

* terminator unless the {@link #DOTALL} flag is specified.

* By default, the regular expressions <tt>^</tt> and <tt>$</tt> ignore

* line terminators and only match at the beginning and the end, respectively,

* of the entire input sequence. If {@link #MULTILINE} mode is activated then

* <tt>^</tt> matches at the beginning of input and after any line terminator

* except at the end of input. When in {@link #MULTILINE} mode <tt>$</tt>

* matches just before a line terminator or the end of the input sequence.

* <h3><a name="cg">Groups and capturing</a></h3>

* <h4><a name="gnumber">Group number</a></h4>

* Capturing groups are numbered by counting their opening parentheses from

* left to right. In the expression <tt>((A)(B(C)))</tt>, for example, there

* are four such groups:

* <blockquote><table cellpadding=1 cellspacing=0 summary="Capturing group numberings">

* <tr><th>1    </th>

* <td><tt>((A)(B(C)))</tt></td></tr>

* <tr><th>2    </th>

* <td><tt>(A)</tt></td></tr>

* <tr><th>3    </th>

* <td><tt>(B(C))</tt></td></tr>

* <tr><th>4    </th>

* <td><tt>(C)</tt></td></tr>

* </table></blockquote>

* Group zero always stands for the entire expression.

* Capturing groups are so named because, during a match, each subsequence

* of the input sequence that matches such a group is saved. The captured

* subsequence may be used later in the expression, via a back reference, and

* may also be retrieved from the matcher once the match operation is complete.

* <h4><a name="groupname">Group name</a></h4>

* A capturing group can also be assigned a "name", a <tt>named-capturing group</tt>,

* and then be back-referenced later by the "name". Group names are composed of

* the following characters. The first character must be a <tt>letter</tt>.

* <ul>

* <li> The uppercase letters <tt>'A'</tt> through <tt>'Z'</tt>

* (<tt>'\u0041'</tt> through <tt>'\u005a'</tt>),

* <li> The lowercase letters <tt>'a'</tt> through <tt>'z'</tt>

* (<tt>'\u0061'</tt> through <tt>'\u007a'</tt>),

* <li> The digits <tt>'0'</tt> through <tt>'9'</tt>

* (<tt>'\u0030'</tt> through <tt>'\u0039'</tt>),

* </ul>

* A <tt>named-capturing group</tt> is still numbered as described in

* <a href="#gnumber">Group number</a>.

* The captured input associated with a group is always the subsequence

* that the group most recently matched. If a group is evaluated a second time

* because of quantification then its previously-captured value, if any, will

* be retained if the second evaluation fails. Matching the string

* <tt>"aba"</tt> against the expression <tt>(a(b)?)+</tt>, for example, leaves

* group two set to <tt>"b"</tt>. All captured input is discarded at the

* beginning of each match.

* Groups beginning with <tt>(?</tt> are either pure, non-capturing groups

* that do not capture text and do not count towards the group total, or

* named-capturing group.

* <h3> Unicode support </h3>

* This class is in conformance with Level 1 of <a

* href="http://www.unicode.org/reports/tr18/">Unicode Technical

* Standard #18: Unicode Regular Expression</a>, plus RL2.1

* Canonical Equivalents.

*

* Unicode escape sequences such as <tt>\u2014</tt> in Java source code

* are processed as described in section 3.3 of

* <cite>The Java™ Language Specification</cite>.

* Such escape sequences are also implemented directly by the regular-expression

* parser so that Unicode escapes can be used in expressions that are read from

* files or from the keyboard. Thus the strings <tt>"\u2014"</tt> and

* <tt>"\\u2014"</tt>, while not equal, compile into the same pattern, which

* matches the character with hexadecimal value <tt>0x2014</tt>.

*

* A Unicode character can also be represented in a regular-expression by

* using its Hex notation(hexadecimal code point value) directly as described in construct

* <tt>\x{...}</tt>, for example a supplementary character U+2011F

* can be specified as <tt>\x{2011F}</tt>, instead of two consecutive

* Unicode escape sequences of the surrogate pair

* <tt>\uD840</tt><tt>\uDD1F</tt>.

*

* Unicode scripts, blocks, categories and binary properties are written with

* the <tt>\p</tt> and <tt>\P</tt> constructs as in Perl.

* <tt>\p{</tt>prop<tt>}</tt> matches if

* the input has the property prop, while <tt>\P{</tt>prop<tt>}</tt>

* does not match if the input has that property.

*

* Scripts, blocks, categories and binary properties can be used both inside

* and outside of a character class.

*

* <a name="usc">Scripts</a> are specified either with the prefix {@code Is}, as in

* {@code IsHiragana}, or by using the {@code script} keyword (or its short

* form {@code sc})as in {@code script=Hiragana} or {@code sc=Hiragana}.

*

* The script names supported by <code>Pattern</code> are the valid script names

* accepted and defined by

* {@link java.lang.Character.UnicodeScript#forName(String) UnicodeScript.forName}.

*

* <a name="ubc">Blocks</a> are specified with the prefix {@code In}, as in

* {@code InMongolian}, or by using the keyword {@code block} (or its short

* form {@code blk}) as in {@code block=Mongolian} or {@code blk=Mongolian}.

*

* The block names supported by <code>Pattern</code> are the valid block names

* accepted and defined by

* {@link java.lang.Character.UnicodeBlock#forName(String) UnicodeBlock.forName}.

*

* <a name="ucc">Categories</a> may be specified with the optional prefix {@code Is}:

* Both {@code \p{L}} and {@code \p{IsL}} denote the category of Unicode

* letters. Same as scripts and blocks, categories can also be specified

* by using the keyword {@code general_category} (or its short form

* {@code gc}) as in {@code general_category=Lu} or {@code gc=Lu}.

*

* The supported categories are those of

* <a href="http://www.unicode.org/unicode/standard/standard.html">

* The Unicode Standard</a> in the version specified by the

* {@link java.lang.Character Character} class. The category names are those

* defined in the Standard, both normative and informative.

*

* <a name="ubpc">Binary properties</a> are specified with the prefix {@code Is}, as in

* {@code IsAlphabetic}. The supported binary properties by <code>Pattern</code>

* are

* <ul>

* <li> Alphabetic

* <li> Ideographic

* <li> Letter

* <li> Lowercase

* <li> Uppercase

* <li> Titlecase

* <li> Punctuation

* <Li> Control

* <li> White_Space

* <li> Digit

* <li> Hex_Digit

* <li> Join_Control

* <li> Noncharacter_Code_Point

* <li> Assigned

* </ul>

*

* The following Predefined Character classes and POSIX character classes

* are in conformance with the recommendation of Annex C: Compatibility Properties

* of <a href="http://www.unicode.org/reports/tr18/">Unicode Regular Expression

* </a>, when {@link #UNICODE_CHARACTER_CLASS} flag is specified.

* <table border="0" cellpadding="1" cellspacing="0"

* summary="predefined and posix character classes in Unicode mode">

* <tr align="left">

* <th align="left" id="predef_classes">Classes</th>

* <th align="left" id="predef_matches">Matches</th>

*</tr>

* <tr><td><tt>\p{Lower}</tt></td>

* <td>A lowercase character:<tt>\p{IsLowercase}</tt></td></tr>

* <tr><td><tt>\p{Upper}</tt></td>

* <td>An uppercase character:<tt>\p{IsUppercase}</tt></td></tr>

* <tr><td><tt>\p{ASCII}</tt></td>

* <td>All ASCII:<tt>[\x00-\x7F]</tt></td></tr>

* <tr><td><tt>\p{Alpha}</tt></td>

* <td>An alphabetic character:<tt>\p{IsAlphabetic}</tt></td></tr>

* <tr><td><tt>\p{Digit}</tt></td>

* <td>A decimal digit character:<tt>p{IsDigit}</tt></td></tr>

* <tr><td><tt>\p{Alnum}</tt></td>

* <td>An alphanumeric character:<tt>[\p{IsAlphabetic}\p{IsDigit}]</tt></td></tr>

* <tr><td><tt>\p{Punct}</tt></td>

* <td>A punctuation character:<tt>p{IsPunctuation}</tt></td></tr>

* <tr><td><tt>\p{Graph}</tt></td>

* <td>A visible character: <tt>[^\p{IsWhite_Space}\p{gc=Cc}\p{gc=Cs}\p{gc=Cn}]</tt></td></tr>

* <tr><td><tt>\p{Print}</tt></td>

* <td>A printable character: {@code [\p{Graph}\p{Blank}&&[^\p{Cntrl}]]}</td></tr>

* <tr><td><tt>\p{Blank}</tt></td>

* <td>A space or a tab: {@code [\p{IsWhite_Space}&&[^\p{gc=Zl}\p{gc=Zp}\x0a\x0b\x0c\x0d\x85]]}</td></tr>

* <tr><td><tt>\p{Cntrl}</tt></td>

* <td>A control character: <tt>\p{gc=Cc}</tt></td></tr>

* <tr><td><tt>\p{XDigit}</tt></td>

* <td>A hexadecimal digit: <tt>[\p{gc=Nd}\p{IsHex_Digit}]</tt></td></tr>

* <tr><td><tt>\p{Space}</tt></td>

* <td>A whitespace character:<tt>\p{IsWhite_Space}</tt></td></tr>

* <tr><td><tt>\d</tt></td>

* <td>A digit: <tt>\p{IsDigit}</tt></td></tr>

* <tr><td><tt>\D</tt></td>

* <td>A non-digit: <tt>[^\d]</tt></td></tr>

* <tr><td><tt>\s</tt></td>

* <td>A whitespace character: <tt>\p{IsWhite_Space}</tt></td></tr>

* <tr><td><tt>\S</tt></td>

* <td>A non-whitespace character: <tt>[^\s]</tt></td></tr>

* <tr><td><tt>\w</tt></td>

* <td>A word character: <tt>[\p{Alpha}\p{gc=Mn}\p{gc=Me}\p{gc=Mc}\p{Digit}\p{gc=Pc}\p{IsJoin_Control}]</tt></td></tr>

* <tr><td><tt>\W</tt></td>

* <td>A non-word character: <tt>[^\w]</tt></td></tr>

* </table>

*

* <a name="jcc">

* Categories that behave like the java.lang.Character

* boolean ismethodname methods (except for the deprecated ones) are

* available through the same <tt>\p{</tt>prop<tt>}</tt> syntax where

* the specified property has the name <tt>javamethodname</tt></a>.

* <h3> Comparison to Perl 5 </h3>

* The <code>Pattern</code> engine performs traditional NFA-based matching

* with ordered alternation as occurs in Perl 5.

* Perl constructs not supported by this class:

* <ul>

* <li> Predefined character classes (Unicode character)

* <tt>\X    </tt>Match Unicode

* <a href="http://www.unicode.org/reports/tr18/#Default_Grapheme_Clusters">

* extended grapheme cluster</a>

* </li>

* <li> The backreference constructs, <tt>\g{</tt>n<tt>}</tt> for

* the nth<a href="#cg">capturing group</a> and

* <tt>\g{</tt>name<tt>}</tt> for

* <a href="#groupname">named-capturing group</a>.

* </li>

* <li> The named character construct, <tt>\N{</tt>name<tt>}</tt>

* for a Unicode character by its name.

* </li>

* <li> The conditional constructs

* <tt>(?(</tt>condition<tt>)</tt>X<tt>)</tt> and

* <tt>(?(</tt>condition<tt>)</tt>X<tt>|</tt>Y<tt>)</tt>,

* </li>

* <li> The embedded code constructs <tt>(?{</tt>code<tt>})</tt>

* and <tt>(??{</tt>code<tt>})</tt>,</li>

* <li> The embedded comment syntax <tt>(?#comment)</tt>, and </li>

* <li> The preprocessing operations <tt>\l</tt> <tt>\u</tt>,

* <tt>\L</tt>, and <tt>\U</tt>. </li>

* </ul>

* Constructs supported by this class but not by Perl:

* <ul>

* <li> Character-class union and intersection as described

* <a href="#cc">above</a>.</li>

* </ul>

* Notable differences from Perl:

* <ul>

* <li> In Perl, <tt>\1</tt> through <tt>\9</tt> are always interpreted

* as back references; a backslash-escaped number greater than <tt>9</tt> is

* treated as a back reference if at least that many subexpressions exist,

* otherwise it is interpreted, if possible, as an octal escape. In this

* class octal escapes must always begin with a zero. In this class,

* <tt>\1</tt> through <tt>\9</tt> are always interpreted as back

* references, and a larger number is accepted as a back reference if at

* least that many subexpressions exist at that point in the regular

* expression, otherwise the parser will drop digits until the number is

* smaller or equal to the existing number of groups or it is one digit.

* </li>

* <li> Perl uses the <tt>g</tt> flag to request a match that resumes

* where the last match left off. This functionality is provided implicitly

* by the {@link Matcher} class: Repeated invocations of the {@link

* Matcher#find find} method will resume where the last match left off,

* unless the matcher is reset. </li>

* <li> In Perl, embedded flags at the top level of an expression affect

* the whole expression. In this class, embedded flags always take effect

* at the point at which they appear, whether they are at the top level or

* within a group; in the latter case, flags are restored at the end of the

* group just as in Perl. </li>

* </ul>

* For a more precise description of the behavior of regular expression

* constructs, please see <a href="http://www.oreilly.com/catalog/regex3/">

* Mastering Regular Expressions, 3nd Edition, Jeffrey E. F. Friedl,

* O'Reilly and Associates, 2006.</a>

*

* @see java.lang.String#split(String, int)

* @see java.lang.String#split(String)

* @author Mike McCloskey

* @author Mark Reinhold

* @author JSR-51 Expert Group

* @since 1.4

* @spec JSR-51

public final class Pattern

implements java.io.Serializable

{

/**

* Regular expression modifier values. Instead of being passed as

* arguments, they can also be passed as inline modifiers.

* For example, the following statements have the same effect.

* <pre>

* RegExp r1 = RegExp.compile("abc", Pattern.I|Pattern.M);

* RegExp r2 = RegExp.compile("(?im)abc", 0);

* </pre>

* The flags are duplicated so that the familiar Perl match flag

* names are available.

/**

* Enables Unix lines mode.

* In this mode, only the <tt>'\n'</tt> line terminator is recognized

* in the behavior of <tt>.</tt>, <tt>^</tt>, and <tt>$</tt>.

* Unix lines mode can also be enabled via the embedded flag

* expression <tt>(?d)</tt>.

public static final int UNIX_LINES = 0x01;

/**

* Enables case-insensitive matching.

* By default, case-insensitive matching assumes that only characters

* in the US-ASCII charset are being matched. Unicode-aware

* case-insensitive matching can be enabled by specifying the {@link

* #UNICODE_CASE} flag in conjunction with this flag.

* Case-insensitive matching can also be enabled via the embedded flag

* expression <tt>(?i)</tt>.

* Specifying this flag may impose a slight performance penalty.

public static final int CASE_INSENSITIVE = 0x02;

/**

* Permits whitespace and comments in pattern.

* In this mode, whitespace is ignored, and embedded comments starting

* with <tt>#</tt> are ignored until the end of a line.

* Comments mode can also be enabled via the embedded flag

* expression <tt>(?x)</tt>.

public static final int COMMENTS = 0x04;

/**

* Enables multiline mode.

* In multiline mode the expressions <tt>^</tt> and <tt>$</tt> match

* just after or just before, respectively, a line terminator or the end of

* the input sequence. By default these expressions only match at the

* beginning and the end of the entire input sequence.

* Multiline mode can also be enabled via the embedded flag

* expression <tt>(?m)</tt>.

public static final int MULTILINE = 0x08;

/**

* Enables literal parsing of the pattern.

* When this flag is specified then the input string that specifies

* the pattern is treated as a sequence of literal characters.

* Metacharacters or escape sequences in the input sequence will be

* given no special meaning.

* The flags CASE_INSENSITIVE and UNICODE_CASE retain their impact on

* matching when used in conjunction with this flag. The other flags

* become superfluous.

* There is no embedded flag character for enabling literal parsing.

* @since 1.5

public static final int LITERAL = 0x10;

/**

* Enables dotall mode.

* In dotall mode, the expression <tt>.</tt> matches any character,

* including a line terminator. By default this expression does not match

* line terminators.

* Dotall mode can also be enabled via the embedded flag

* expression <tt>(?s)</tt>. (The <tt>s</tt> is a mnemonic for

* "single-line" mode, which is what this is called in Perl.)

public static final int DOTALL = 0x20;

/**

* Enables Unicode-aware case folding.

* When this flag is specified then case-insensitive matching, when

* enabled by the {@link #CASE_INSENSITIVE} flag, is done in a manner

* consistent with the Unicode Standard. By default, case-insensitive

* matching assumes that only characters in the US-ASCII charset are being

* matched.

* Unicode-aware case folding can also be enabled via the embedded flag

* expression <tt>(?u)</tt>.

* Specifying this flag may impose a performance penalty.

public static final int UNICODE_CASE = 0x40;

/**

* Enables canonical equivalence.

* When this flag is specified then two characters will be considered

* to match if, and only if, their full canonical decompositions match.

* The expression <tt>"a\u030A"</tt>, for example, will match the

* string <tt>"\u00E5"</tt> when this flag is specified. By default,

* matching does not take canonical equivalence into account.

* There is no embedded flag character for enabling canonical

* equivalence.

* Specifying this flag may impose a performance penalty.

public static final int CANON_EQ = 0x80;

/**

* Enables the Unicode version of Predefined character classes and

* POSIX character classes.

* When this flag is specified then the (US-ASCII only)

* Predefined character classes and POSIX character classes

* are in conformance with

* <a href="http://www.unicode.org/reports/tr18/">Unicode Technical

* Standard #18: Unicode Regular Expression</a>

* Annex C: Compatibility Properties.

*

* The UNICODE_CHARACTER_CLASS mode can also be enabled via the embedded

* flag expression <tt>(?U)</tt>.

*

* The flag implies UNICODE_CASE, that is, it enables Unicode-aware case

* folding.

*

* Specifying this flag may impose a performance penalty.

* @since 1.7

public static final int UNICODE_CHARACTER_CLASS = 0x100;

/* Pattern has only two serialized components: The pattern string

* and the flags, which are all that is needed to recompile the pattern

* when it is deserialized.

/** use serialVersionUID from Merlin b59 for interoperability */

private static final long serialVersionUID = 5073258162644648461L;

/**

* The original regular-expression pattern string.

* @serial

private String pattern;

/**

* The original pattern flags.

* @serial

private int flags;

/**

* Boolean indicating this Pattern is compiled; this is necessary in order

* to lazily compile deserialized Patterns.

private transient volatile boolean compiled = false;

/**

* The normalized pattern string.

private transient String normalizedPattern;

/**

* The starting point of state machine for the find operation. This allows

* a match to start anywhere in the input.

transient Node root;

/**

* The root of object tree for a match operation. The pattern is matched

* at the beginning. This may include a find that uses BnM or a First

* node.

transient Node matchRoot;

/**

* Temporary storage used by parsing pattern slice.

transient int[] buffer;

/**

* Map the "name" of the "named capturing group" to its group id

* node.

transient volatile Map<String, Integer> namedGroups;

/**

* Temporary storage used while parsing group references.

transient GroupHead[] groupNodes;

/**

* Temporary null terminated code point array used by pattern compiling.

private transient int[] temp;

/**

* The number of capturing groups in this Pattern. Used by matchers to

* allocate storage needed to perform a match.

transient int capturingGroupCount;

/**

* The local variable count used by parsing tree. Used by matchers to

* allocate storage needed to perform a match.

transient int localCount;

/**

* Index into the pattern string that keeps track of how much has been

View remainder of file in raw view

Provide feedback

Saved searches

Use saved searches to filter your results more quickly

Expand file tree

Search code, repositories, users, issues, pull requests...

FilesExpand file tree

Pattern.java

Latest commit

History

Pattern.java

File metadata and controls

Expand file tree