blob: d89c79b3d93d16c488b4e63fc18cc08b861e2795 [file] [log] [blame]
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +00001//===--- CommentParser.cpp - Doxygen comment parser -----------------------===//
2//
3// The LLVM Compiler Infrastructure
4//
5// This file is distributed under the University of Illinois Open Source
6// License. See LICENSE.TXT for details.
7//
8//===----------------------------------------------------------------------===//
9
10#include "clang/AST/CommentParser.h"
Dmitri Gribenkoaa580812012-08-09 00:03:17 +000011#include "clang/AST/CommentCommandTraits.h"
Chandler Carruth55fc8732012-12-04 09:13:33 +000012#include "clang/AST/CommentDiagnostic.h"
13#include "clang/AST/CommentSema.h"
Dmitri Gribenkobf881442013-02-09 15:16:58 +000014#include "clang/Basic/CharInfo.h"
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +000015#include "clang/Basic/SourceManager.h"
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +000016#include "llvm/Support/ErrorHandling.h"
17
18namespace clang {
19namespace comments {
20
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +000021/// Re-lexes a sequence of tok::text tokens.
22class TextTokenRetokenizer {
23 llvm::BumpPtrAllocator &Allocator;
Dmitri Gribenkodb13f042012-07-24 17:52:18 +000024 Parser &P;
Dmitri Gribenko0c43a922012-07-24 18:23:31 +000025
26 /// This flag is set when there are no more tokens we can fetch from lexer.
27 bool NoMoreInterestingTokens;
28
29 /// Token buffer: tokens we have processed and lookahead.
Dmitri Gribenkodb13f042012-07-24 17:52:18 +000030 SmallVector<Token, 16> Toks;
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +000031
Dmitri Gribenko0c43a922012-07-24 18:23:31 +000032 /// A position in \c Toks.
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +000033 struct Position {
34 unsigned CurToken;
35 const char *BufferStart;
36 const char *BufferEnd;
37 const char *BufferPtr;
38 SourceLocation BufferStartLoc;
39 };
40
41 /// Current position in Toks.
42 Position Pos;
43
44 bool isEnd() const {
45 return Pos.CurToken >= Toks.size();
46 }
47
48 /// Sets up the buffer pointers to point to current token.
49 void setupBuffer() {
Dmitri Gribenkodb13f042012-07-24 17:52:18 +000050 assert(!isEnd());
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +000051 const Token &Tok = Toks[Pos.CurToken];
52
53 Pos.BufferStart = Tok.getText().begin();
54 Pos.BufferEnd = Tok.getText().end();
55 Pos.BufferPtr = Pos.BufferStart;
56 Pos.BufferStartLoc = Tok.getLocation();
57 }
58
59 SourceLocation getSourceLocation() const {
60 const unsigned CharNo = Pos.BufferPtr - Pos.BufferStart;
61 return Pos.BufferStartLoc.getLocWithOffset(CharNo);
62 }
63
64 char peek() const {
65 assert(!isEnd());
66 assert(Pos.BufferPtr != Pos.BufferEnd);
67 return *Pos.BufferPtr;
68 }
69
70 void consumeChar() {
71 assert(!isEnd());
72 assert(Pos.BufferPtr != Pos.BufferEnd);
73 Pos.BufferPtr++;
74 if (Pos.BufferPtr == Pos.BufferEnd) {
75 Pos.CurToken++;
Dmitri Gribenko0c43a922012-07-24 18:23:31 +000076 if (isEnd() && !addToken())
77 return;
78
79 assert(!isEnd());
80 setupBuffer();
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +000081 }
82 }
83
Dmitri Gribenkodb13f042012-07-24 17:52:18 +000084 /// Add a token.
85 /// Returns true on success, false if there are no interesting tokens to
86 /// fetch from lexer.
87 bool addToken() {
Dmitri Gribenko0c43a922012-07-24 18:23:31 +000088 if (NoMoreInterestingTokens)
Dmitri Gribenkodb13f042012-07-24 17:52:18 +000089 return false;
90
Dmitri Gribenko0c43a922012-07-24 18:23:31 +000091 if (P.Tok.is(tok::newline)) {
92 // If we see a single newline token between text tokens, skip it.
93 Token Newline = P.Tok;
94 P.consumeToken();
95 if (P.Tok.isNot(tok::text)) {
96 P.putBack(Newline);
97 NoMoreInterestingTokens = true;
98 return false;
99 }
100 }
101 if (P.Tok.isNot(tok::text)) {
102 NoMoreInterestingTokens = true;
103 return false;
104 }
105
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000106 Toks.push_back(P.Tok);
107 P.consumeToken();
108 if (Toks.size() == 1)
109 setupBuffer();
110 return true;
111 }
112
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000113 void consumeWhitespace() {
114 while (!isEnd()) {
115 if (isWhitespace(peek()))
116 consumeChar();
117 else
118 break;
119 }
120 }
121
122 void formTokenWithChars(Token &Result,
123 SourceLocation Loc,
124 const char *TokBegin,
125 unsigned TokLength,
126 StringRef Text) {
127 Result.setLocation(Loc);
128 Result.setKind(tok::text);
129 Result.setLength(TokLength);
130#ifndef NDEBUG
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000131 Result.TextPtr = "<UNSET>";
132 Result.IntVal = 7;
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000133#endif
134 Result.setText(Text);
135 }
136
137public:
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000138 TextTokenRetokenizer(llvm::BumpPtrAllocator &Allocator, Parser &P):
Dmitri Gribenko0c43a922012-07-24 18:23:31 +0000139 Allocator(Allocator), P(P), NoMoreInterestingTokens(false) {
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000140 Pos.CurToken = 0;
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000141 addToken();
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000142 }
143
144 /// Extract a word -- sequence of non-whitespace characters.
145 bool lexWord(Token &Tok) {
146 if (isEnd())
147 return false;
148
149 Position SavedPos = Pos;
150
151 consumeWhitespace();
152 SmallString<32> WordText;
153 const char *WordBegin = Pos.BufferPtr;
154 SourceLocation Loc = getSourceLocation();
155 while (!isEnd()) {
156 const char C = peek();
157 if (!isWhitespace(C)) {
158 WordText.push_back(C);
159 consumeChar();
160 } else
161 break;
162 }
163 const unsigned Length = WordText.size();
164 if (Length == 0) {
165 Pos = SavedPos;
166 return false;
167 }
168
169 char *TextPtr = Allocator.Allocate<char>(Length + 1);
170
171 memcpy(TextPtr, WordText.c_str(), Length + 1);
172 StringRef Text = StringRef(TextPtr, Length);
173
Dmitri Gribenkoca57ccd2012-12-19 17:34:55 +0000174 formTokenWithChars(Tok, Loc, WordBegin, Length, Text);
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000175 return true;
176 }
177
178 bool lexDelimitedSeq(Token &Tok, char OpenDelim, char CloseDelim) {
179 if (isEnd())
180 return false;
181
182 Position SavedPos = Pos;
183
184 consumeWhitespace();
185 SmallString<32> WordText;
186 const char *WordBegin = Pos.BufferPtr;
187 SourceLocation Loc = getSourceLocation();
188 bool Error = false;
189 if (!isEnd()) {
190 const char C = peek();
191 if (C == OpenDelim) {
192 WordText.push_back(C);
193 consumeChar();
194 } else
195 Error = true;
196 }
197 char C = '\0';
198 while (!Error && !isEnd()) {
199 C = peek();
200 WordText.push_back(C);
201 consumeChar();
202 if (C == CloseDelim)
203 break;
204 }
205 if (!Error && C != CloseDelim)
206 Error = true;
207
208 if (Error) {
209 Pos = SavedPos;
210 return false;
211 }
212
213 const unsigned Length = WordText.size();
214 char *TextPtr = Allocator.Allocate<char>(Length + 1);
215
216 memcpy(TextPtr, WordText.c_str(), Length + 1);
217 StringRef Text = StringRef(TextPtr, Length);
218
219 formTokenWithChars(Tok, Loc, WordBegin,
220 Pos.BufferPtr - WordBegin, Text);
221 return true;
222 }
223
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000224 /// Put back tokens that we didn't consume.
225 void putBackLeftoverTokens() {
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000226 if (isEnd())
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000227 return;
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000228
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000229 bool HavePartialTok = false;
230 Token PartialTok;
231 if (Pos.BufferPtr != Pos.BufferStart) {
232 formTokenWithChars(PartialTok, getSourceLocation(),
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000233 Pos.BufferPtr, Pos.BufferEnd - Pos.BufferPtr,
234 StringRef(Pos.BufferPtr,
235 Pos.BufferEnd - Pos.BufferPtr));
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000236 HavePartialTok = true;
237 Pos.CurToken++;
238 }
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000239
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000240 P.putBack(llvm::makeArrayRef(Toks.begin() + Pos.CurToken, Toks.end()));
241 Pos.CurToken = Toks.size();
242
243 if (HavePartialTok)
244 P.putBack(PartialTok);
Dmitri Gribenkoc4b0f9b2012-07-24 17:43:18 +0000245 }
246};
247
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000248Parser::Parser(Lexer &L, Sema &S, llvm::BumpPtrAllocator &Allocator,
Dmitri Gribenkoaa580812012-08-09 00:03:17 +0000249 const SourceManager &SourceMgr, DiagnosticsEngine &Diags,
250 const CommandTraits &Traits):
251 L(L), S(S), Allocator(Allocator), SourceMgr(SourceMgr), Diags(Diags),
252 Traits(Traits) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000253 consumeToken();
254}
255
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000256void Parser::parseParamCommandArgs(ParamCommandComment *PC,
257 TextTokenRetokenizer &Retokenizer) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000258 Token Arg;
259 // Check if argument looks like direction specification: [dir]
260 // e.g., [in], [out], [in,out]
261 if (Retokenizer.lexDelimitedSeq(Arg, '[', ']'))
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000262 S.actOnParamCommandDirectionArg(PC,
263 Arg.getLocation(),
264 Arg.getEndLocation(),
265 Arg.getText());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000266
267 if (Retokenizer.lexWord(Arg))
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000268 S.actOnParamCommandParamNameArg(PC,
269 Arg.getLocation(),
270 Arg.getEndLocation(),
271 Arg.getText());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000272}
273
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000274void Parser::parseTParamCommandArgs(TParamCommandComment *TPC,
275 TextTokenRetokenizer &Retokenizer) {
Dmitri Gribenko96b09862012-07-31 22:37:06 +0000276 Token Arg;
277 if (Retokenizer.lexWord(Arg))
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000278 S.actOnTParamCommandParamNameArg(TPC,
279 Arg.getLocation(),
280 Arg.getEndLocation(),
281 Arg.getText());
Dmitri Gribenko96b09862012-07-31 22:37:06 +0000282}
283
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000284void Parser::parseBlockCommandArgs(BlockCommandComment *BC,
285 TextTokenRetokenizer &Retokenizer,
286 unsigned NumArgs) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000287 typedef BlockCommandComment::Argument Argument;
Dmitri Gribenko814e2192012-07-06 16:41:59 +0000288 Argument *Args =
289 new (Allocator.Allocate<Argument>(NumArgs)) Argument[NumArgs];
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000290 unsigned ParsedArgs = 0;
291 Token Arg;
292 while (ParsedArgs < NumArgs && Retokenizer.lexWord(Arg)) {
293 Args[ParsedArgs] = Argument(SourceRange(Arg.getLocation(),
294 Arg.getEndLocation()),
295 Arg.getText());
296 ParsedArgs++;
297 }
298
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000299 S.actOnBlockCommandArgs(BC, llvm::makeArrayRef(Args, ParsedArgs));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000300}
301
302BlockCommandComment *Parser::parseBlockCommand() {
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000303 assert(Tok.is(tok::backslash_command) || Tok.is(tok::at_command));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000304
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000305 ParamCommandComment *PC = 0;
306 TParamCommandComment *TPC = 0;
307 BlockCommandComment *BC = 0;
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000308 const CommandInfo *Info = Traits.getCommandInfo(Tok.getCommandID());
Dmitri Gribenko808383d2013-03-04 23:06:15 +0000309 CommandMarkerKind CommandMarker =
310 Tok.is(tok::backslash_command) ? CMK_Backslash : CMK_At;
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000311 if (Info->IsParamCommand) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000312 PC = S.actOnParamCommandStart(Tok.getLocation(),
313 Tok.getEndLocation(),
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000314 Tok.getCommandID(),
Dmitri Gribenko808383d2013-03-04 23:06:15 +0000315 CommandMarker);
Dmitri Gribenkoeb34db72012-12-19 17:17:09 +0000316 } else if (Info->IsTParamCommand) {
Dmitri Gribenko96b09862012-07-31 22:37:06 +0000317 TPC = S.actOnTParamCommandStart(Tok.getLocation(),
318 Tok.getEndLocation(),
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000319 Tok.getCommandID(),
Dmitri Gribenko808383d2013-03-04 23:06:15 +0000320 CommandMarker);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000321 } else {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000322 BC = S.actOnBlockCommandStart(Tok.getLocation(),
323 Tok.getEndLocation(),
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000324 Tok.getCommandID(),
Dmitri Gribenko808383d2013-03-04 23:06:15 +0000325 CommandMarker);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000326 }
327 consumeToken();
328
Dmitri Gribenko10442562013-01-26 00:36:14 +0000329 if (isTokBlockCommand()) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000330 // Block command ahead. We can't nest block commands, so pretend that this
331 // command has an empty argument.
Dmitri Gribenko55431692013-05-05 00:41:58 +0000332 ParagraphComment *Paragraph = S.actOnParagraphComment(None);
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000333 if (PC) {
Dmitri Gribenko8a903932012-08-06 23:48:44 +0000334 S.actOnParamCommandFinish(PC, Paragraph);
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000335 return PC;
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000336 } else if (TPC) {
Dmitri Gribenko8a903932012-08-06 23:48:44 +0000337 S.actOnTParamCommandFinish(TPC, Paragraph);
338 return TPC;
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000339 } else {
340 S.actOnBlockCommandFinish(BC, Paragraph);
341 return BC;
342 }
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000343 }
344
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000345 if (PC || TPC || Info->NumArgs > 0) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000346 // In order to parse command arguments we need to retokenize a few
347 // following text tokens.
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000348 TextTokenRetokenizer Retokenizer(Allocator, *this);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000349
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000350 if (PC)
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000351 parseParamCommandArgs(PC, Retokenizer);
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000352 else if (TPC)
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000353 parseTParamCommandArgs(TPC, Retokenizer);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000354 else
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000355 parseBlockCommandArgs(BC, Retokenizer, Info->NumArgs);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000356
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000357 Retokenizer.putBackLeftoverTokens();
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000358 }
359
Dmitri Gribenko10442562013-01-26 00:36:14 +0000360 // If there's a block command ahead, we will attach an empty paragraph to
361 // this command.
362 bool EmptyParagraph = false;
363 if (isTokBlockCommand())
364 EmptyParagraph = true;
365 else if (Tok.is(tok::newline)) {
366 Token PrevTok = Tok;
367 consumeToken();
368 EmptyParagraph = isTokBlockCommand();
369 putBack(PrevTok);
370 }
371
372 ParagraphComment *Paragraph;
373 if (EmptyParagraph)
Dmitri Gribenko55431692013-05-05 00:41:58 +0000374 Paragraph = S.actOnParagraphComment(None);
Dmitri Gribenko10442562013-01-26 00:36:14 +0000375 else {
376 BlockContentComment *Block = parseParagraphOrBlockCommand();
377 // Since we have checked for a block command, we should have parsed a
378 // paragraph.
379 Paragraph = cast<ParagraphComment>(Block);
380 }
381
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000382 if (PC) {
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000383 S.actOnParamCommandFinish(PC, Paragraph);
384 return PC;
Dmitri Gribenko3a589122013-04-18 20:50:35 +0000385 } else if (TPC) {
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000386 S.actOnTParamCommandFinish(TPC, Paragraph);
387 return TPC;
388 } else {
389 S.actOnBlockCommandFinish(BC, Paragraph);
390 return BC;
391 }
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000392}
393
394InlineCommandComment *Parser::parseInlineCommand() {
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000395 assert(Tok.is(tok::backslash_command) || Tok.is(tok::at_command));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000396
397 const Token CommandTok = Tok;
398 consumeToken();
399
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000400 TextTokenRetokenizer Retokenizer(Allocator, *this);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000401
402 Token ArgTok;
403 bool ArgTokValid = Retokenizer.lexWord(ArgTok);
404
405 InlineCommandComment *IC;
406 if (ArgTokValid) {
407 IC = S.actOnInlineCommand(CommandTok.getLocation(),
408 CommandTok.getEndLocation(),
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000409 CommandTok.getCommandID(),
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000410 ArgTok.getLocation(),
411 ArgTok.getEndLocation(),
412 ArgTok.getText());
413 } else {
414 IC = S.actOnInlineCommand(CommandTok.getLocation(),
415 CommandTok.getEndLocation(),
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000416 CommandTok.getCommandID());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000417 }
418
Dmitri Gribenkodb13f042012-07-24 17:52:18 +0000419 Retokenizer.putBackLeftoverTokens();
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000420
421 return IC;
422}
423
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000424HTMLStartTagComment *Parser::parseHTMLStartTag() {
425 assert(Tok.is(tok::html_start_tag));
426 HTMLStartTagComment *HST =
427 S.actOnHTMLStartTagStart(Tok.getLocation(),
428 Tok.getHTMLTagStartName());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000429 consumeToken();
430
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000431 SmallVector<HTMLStartTagComment::Attribute, 2> Attrs;
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000432 while (true) {
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000433 switch (Tok.getKind()) {
434 case tok::html_ident: {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000435 Token Ident = Tok;
436 consumeToken();
437 if (Tok.isNot(tok::html_equals)) {
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000438 Attrs.push_back(HTMLStartTagComment::Attribute(Ident.getLocation(),
439 Ident.getHTMLIdent()));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000440 continue;
441 }
442 Token Equals = Tok;
443 consumeToken();
444 if (Tok.isNot(tok::html_quoted_string)) {
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000445 Diag(Tok.getLocation(),
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000446 diag::warn_doc_html_start_tag_expected_quoted_string)
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000447 << SourceRange(Equals.getLocation());
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000448 Attrs.push_back(HTMLStartTagComment::Attribute(Ident.getLocation(),
449 Ident.getHTMLIdent()));
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000450 while (Tok.is(tok::html_equals) ||
451 Tok.is(tok::html_quoted_string))
452 consumeToken();
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000453 continue;
454 }
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000455 Attrs.push_back(HTMLStartTagComment::Attribute(
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000456 Ident.getLocation(),
457 Ident.getHTMLIdent(),
458 Equals.getLocation(),
459 SourceRange(Tok.getLocation(),
460 Tok.getEndLocation()),
461 Tok.getHTMLQuotedString()));
462 consumeToken();
463 continue;
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000464 }
465
466 case tok::html_greater:
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000467 S.actOnHTMLStartTagFinish(HST,
468 S.copyArray(llvm::makeArrayRef(Attrs)),
469 Tok.getLocation(),
470 /* IsSelfClosing = */ false);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000471 consumeToken();
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000472 return HST;
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000473
474 case tok::html_slash_greater:
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000475 S.actOnHTMLStartTagFinish(HST,
476 S.copyArray(llvm::makeArrayRef(Attrs)),
477 Tok.getLocation(),
478 /* IsSelfClosing = */ true);
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000479 consumeToken();
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000480 return HST;
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000481
482 case tok::html_equals:
483 case tok::html_quoted_string:
484 Diag(Tok.getLocation(),
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000485 diag::warn_doc_html_start_tag_expected_ident_or_greater);
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000486 while (Tok.is(tok::html_equals) ||
487 Tok.is(tok::html_quoted_string))
488 consumeToken();
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000489 if (Tok.is(tok::html_ident) ||
490 Tok.is(tok::html_greater) ||
491 Tok.is(tok::html_slash_greater))
492 continue;
493
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000494 S.actOnHTMLStartTagFinish(HST,
495 S.copyArray(llvm::makeArrayRef(Attrs)),
496 SourceLocation(),
497 /* IsSelfClosing = */ false);
498 return HST;
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000499
500 default:
501 // Not a token from an HTML start tag. Thus HTML tag prematurely ended.
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000502 S.actOnHTMLStartTagFinish(HST,
503 S.copyArray(llvm::makeArrayRef(Attrs)),
504 SourceLocation(),
505 /* IsSelfClosing = */ false);
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000506 bool StartLineInvalid;
507 const unsigned StartLine = SourceMgr.getPresumedLineNumber(
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000508 HST->getLocation(),
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000509 &StartLineInvalid);
510 bool EndLineInvalid;
511 const unsigned EndLine = SourceMgr.getPresumedLineNumber(
512 Tok.getLocation(),
513 &EndLineInvalid);
514 if (StartLineInvalid || EndLineInvalid || StartLine == EndLine)
515 Diag(Tok.getLocation(),
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000516 diag::warn_doc_html_start_tag_expected_ident_or_greater)
517 << HST->getSourceRange();
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000518 else {
519 Diag(Tok.getLocation(),
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000520 diag::warn_doc_html_start_tag_expected_ident_or_greater);
521 Diag(HST->getLocation(), diag::note_doc_html_tag_started_here)
522 << HST->getSourceRange();
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000523 }
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000524 return HST;
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000525 }
526 }
527}
528
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000529HTMLEndTagComment *Parser::parseHTMLEndTag() {
530 assert(Tok.is(tok::html_end_tag));
531 Token TokEndTag = Tok;
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000532 consumeToken();
533 SourceLocation Loc;
534 if (Tok.is(tok::html_greater)) {
535 Loc = Tok.getLocation();
536 consumeToken();
537 }
538
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000539 return S.actOnHTMLEndTag(TokEndTag.getLocation(),
540 Loc,
541 TokEndTag.getHTMLTagEndName());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000542}
543
544BlockContentComment *Parser::parseParagraphOrBlockCommand() {
545 SmallVector<InlineContentComment *, 8> Content;
546
547 while (true) {
548 switch (Tok.getKind()) {
549 case tok::verbatim_block_begin:
550 case tok::verbatim_line_name:
551 case tok::eof:
552 assert(Content.size() != 0);
553 break; // Block content or EOF ahead, finish this parapgaph.
554
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000555 case tok::unknown_command:
556 Content.push_back(S.actOnUnknownCommand(Tok.getLocation(),
557 Tok.getEndLocation(),
558 Tok.getUnknownCommandName()));
559 consumeToken();
560 continue;
561
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000562 case tok::backslash_command:
563 case tok::at_command: {
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000564 const CommandInfo *Info = Traits.getCommandInfo(Tok.getCommandID());
565 if (Info->IsBlockCommand) {
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000566 if (Content.size() == 0)
567 return parseBlockCommand();
568 break; // Block command ahead, finish this parapgaph.
569 }
Dmitri Gribenko36cbbe92012-11-18 00:30:31 +0000570 if (Info->IsVerbatimBlockEndCommand) {
571 Diag(Tok.getLocation(),
572 diag::warn_verbatim_block_end_without_start)
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000573 << Tok.is(tok::at_command)
Dmitri Gribenko36cbbe92012-11-18 00:30:31 +0000574 << Info->Name
575 << SourceRange(Tok.getLocation(), Tok.getEndLocation());
576 consumeToken();
577 continue;
578 }
Dmitri Gribenkob0b8a962012-09-11 19:22:03 +0000579 if (Info->IsUnknownCommand) {
580 Content.push_back(S.actOnUnknownCommand(Tok.getLocation(),
581 Tok.getEndLocation(),
582 Info->getID()));
583 consumeToken();
584 continue;
585 }
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000586 assert(Info->IsInlineCommand);
587 Content.push_back(parseInlineCommand());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000588 continue;
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000589 }
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000590
591 case tok::newline: {
592 consumeToken();
593 if (Tok.is(tok::newline) || Tok.is(tok::eof)) {
594 consumeToken();
595 break; // Two newlines -- end of paragraph.
596 }
597 if (Content.size() > 0)
598 Content.back()->addTrailingNewline();
599 continue;
600 }
601
602 // Don't deal with HTML tag soup now.
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000603 case tok::html_start_tag:
604 Content.push_back(parseHTMLStartTag());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000605 continue;
606
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000607 case tok::html_end_tag:
608 Content.push_back(parseHTMLEndTag());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000609 continue;
610
611 case tok::text:
612 Content.push_back(S.actOnText(Tok.getLocation(),
613 Tok.getEndLocation(),
614 Tok.getText()));
615 consumeToken();
616 continue;
617
618 case tok::verbatim_block_line:
619 case tok::verbatim_block_end:
620 case tok::verbatim_line_text:
621 case tok::html_ident:
622 case tok::html_equals:
623 case tok::html_quoted_string:
624 case tok::html_greater:
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000625 case tok::html_slash_greater:
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000626 llvm_unreachable("should not see this token");
627 }
628 break;
629 }
630
Dmitri Gribenko96b09862012-07-31 22:37:06 +0000631 return S.actOnParagraphComment(S.copyArray(llvm::makeArrayRef(Content)));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000632}
633
634VerbatimBlockComment *Parser::parseVerbatimBlock() {
635 assert(Tok.is(tok::verbatim_block_begin));
636
637 VerbatimBlockComment *VB =
638 S.actOnVerbatimBlockStart(Tok.getLocation(),
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000639 Tok.getVerbatimBlockID());
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000640 consumeToken();
641
642 // Don't create an empty line if verbatim opening command is followed
643 // by a newline.
644 if (Tok.is(tok::newline))
645 consumeToken();
646
647 SmallVector<VerbatimBlockLineComment *, 8> Lines;
648 while (Tok.is(tok::verbatim_block_line) ||
649 Tok.is(tok::newline)) {
650 VerbatimBlockLineComment *Line;
651 if (Tok.is(tok::verbatim_block_line)) {
652 Line = S.actOnVerbatimBlockLine(Tok.getLocation(),
653 Tok.getVerbatimBlockText());
654 consumeToken();
655 if (Tok.is(tok::newline)) {
656 consumeToken();
657 }
658 } else {
659 // Empty line, just a tok::newline.
Dmitri Gribenko94572c32012-07-18 21:27:38 +0000660 Line = S.actOnVerbatimBlockLine(Tok.getLocation(), "");
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000661 consumeToken();
662 }
663 Lines.push_back(Line);
664 }
665
Dmitri Gribenko9f08f492012-07-20 20:18:53 +0000666 if (Tok.is(tok::verbatim_block_end)) {
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000667 const CommandInfo *Info = Traits.getCommandInfo(Tok.getVerbatimBlockID());
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000668 S.actOnVerbatimBlockFinish(VB, Tok.getLocation(),
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000669 Info->Name,
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000670 S.copyArray(llvm::makeArrayRef(Lines)));
Dmitri Gribenko9f08f492012-07-20 20:18:53 +0000671 consumeToken();
672 } else {
673 // Unterminated \\verbatim block
Dmitri Gribenko7d9b5112012-08-06 19:03:12 +0000674 S.actOnVerbatimBlockFinish(VB, SourceLocation(), "",
675 S.copyArray(llvm::makeArrayRef(Lines)));
Dmitri Gribenko9f08f492012-07-20 20:18:53 +0000676 }
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000677
678 return VB;
679}
680
681VerbatimLineComment *Parser::parseVerbatimLine() {
682 assert(Tok.is(tok::verbatim_line_name));
683
684 Token NameTok = Tok;
685 consumeToken();
686
687 SourceLocation TextBegin;
688 StringRef Text;
689 // Next token might not be a tok::verbatim_line_text if verbatim line
690 // starting command comes just before a newline or comment end.
691 if (Tok.is(tok::verbatim_line_text)) {
692 TextBegin = Tok.getLocation();
693 Text = Tok.getVerbatimLineText();
694 } else {
695 TextBegin = NameTok.getEndLocation();
696 Text = "";
697 }
698
699 VerbatimLineComment *VL = S.actOnVerbatimLine(NameTok.getLocation(),
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000700 NameTok.getVerbatimLineID(),
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000701 TextBegin,
702 Text);
703 consumeToken();
704 return VL;
705}
706
707BlockContentComment *Parser::parseBlockContent() {
708 switch (Tok.getKind()) {
709 case tok::text:
Dmitri Gribenkoe4330a32012-09-10 20:32:42 +0000710 case tok::unknown_command:
Fariborz Jahanian8536fa12013-03-02 02:39:57 +0000711 case tok::backslash_command:
712 case tok::at_command:
Dmitri Gribenko3f38bf22012-07-13 00:44:24 +0000713 case tok::html_start_tag:
714 case tok::html_end_tag:
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000715 return parseParagraphOrBlockCommand();
716
717 case tok::verbatim_block_begin:
718 return parseVerbatimBlock();
719
720 case tok::verbatim_line_name:
721 return parseVerbatimLine();
722
723 case tok::eof:
724 case tok::newline:
725 case tok::verbatim_block_line:
726 case tok::verbatim_block_end:
727 case tok::verbatim_line_text:
728 case tok::html_ident:
729 case tok::html_equals:
730 case tok::html_quoted_string:
731 case tok::html_greater:
Dmitri Gribenkoa5ef44f2012-07-11 21:38:39 +0000732 case tok::html_slash_greater:
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000733 llvm_unreachable("should not see this token");
734 }
Matt Beaumont-Gay4d48b5c2012-07-06 21:13:09 +0000735 llvm_unreachable("bogus token kind");
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000736}
737
738FullComment *Parser::parseFullComment() {
739 // Skip newlines at the beginning of the comment.
740 while (Tok.is(tok::newline))
741 consumeToken();
742
743 SmallVector<BlockContentComment *, 8> Blocks;
744 while (Tok.isNot(tok::eof)) {
745 Blocks.push_back(parseBlockContent());
746
747 // Skip extra newlines after paragraph end.
748 while (Tok.is(tok::newline))
749 consumeToken();
750 }
Dmitri Gribenko96b09862012-07-31 22:37:06 +0000751 return S.actOnFullComment(S.copyArray(llvm::makeArrayRef(Blocks)));
Dmitri Gribenko8d3ba232012-07-06 00:28:32 +0000752}
753
754} // end namespace comments
755} // end namespace clang