blob: 4021bd5d92847e0607f494e2081d89bbf5834848 [file] [log] [blame]
Kate Stoneb9c1b512016-09-06 20:57:50 +00001//===-- StringPrinter.cpp ----------------------------------------*- C++
2//-*-===//
Enrico Granataca6c8ee2014-10-30 01:45:39 +00003//
4// The LLVM Compiler Infrastructure
5//
6// This file is distributed under the University of Illinois Open Source
7// License. See LICENSE.TXT for details.
8//
9//===----------------------------------------------------------------------===//
10
11#include "lldb/DataFormatters/StringPrinter.h"
12
Enrico Granataebdc1ac2014-11-05 21:20:48 +000013#include "lldb/Core/Debugger.h"
Enrico Granataca6c8ee2014-10-30 01:45:39 +000014#include "lldb/Core/Error.h"
Enrico Granataebdc1ac2014-11-05 21:20:48 +000015#include "lldb/Core/ValueObject.h"
Enrico Granataac494532015-09-09 22:30:24 +000016#include "lldb/Target/Language.h"
Enrico Granataca6c8ee2014-10-30 01:45:39 +000017#include "lldb/Target/Process.h"
18#include "lldb/Target/Target.h"
19
20#include "llvm/Support/ConvertUTF.h"
21
Enrico Granataca6c8ee2014-10-30 01:45:39 +000022#include <ctype.h>
Enrico Granataca6c8ee2014-10-30 01:45:39 +000023#include <locale>
24
25using namespace lldb;
26using namespace lldb_private;
27using namespace lldb_private::formatters;
28
Kate Stoneb9c1b512016-09-06 20:57:50 +000029// we define this for all values of type but only implement it for those we care
30// about
Enrico Granataca6c8ee2014-10-30 01:45:39 +000031// that's good because we get linker errors for any unsupported type
Enrico Granataac494532015-09-09 22:30:24 +000032template <lldb_private::formatters::StringPrinter::StringElementType type>
Enrico Granataad650a12015-09-09 20:59:49 +000033static StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +000034GetPrintableImpl(uint8_t *buffer, uint8_t *buffer_end, uint8_t *&next);
Enrico Granataca6c8ee2014-10-30 01:45:39 +000035
36// mimic isprint() for Unicode codepoints
Kate Stoneb9c1b512016-09-06 20:57:50 +000037static bool isprint(char32_t codepoint) {
38 if (codepoint <= 0x1F || codepoint == 0x7F) // C0
39 {
40 return false;
41 }
42 if (codepoint >= 0x80 && codepoint <= 0x9F) // C1
43 {
44 return false;
45 }
46 if (codepoint == 0x2028 || codepoint == 0x2029) // line/paragraph separators
47 {
48 return false;
49 }
50 if (codepoint == 0x200E || codepoint == 0x200F ||
51 (codepoint >= 0x202A &&
52 codepoint <= 0x202E)) // bidirectional text control
53 {
54 return false;
55 }
56 if (codepoint >= 0xFFF9 &&
57 codepoint <= 0xFFFF) // interlinears and generally specials
58 {
59 return false;
60 }
61 return true;
Enrico Granataca6c8ee2014-10-30 01:45:39 +000062}
63
64template <>
Enrico Granataad650a12015-09-09 20:59:49 +000065StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +000066GetPrintableImpl<StringPrinter::StringElementType::ASCII>(uint8_t *buffer,
67 uint8_t *buffer_end,
68 uint8_t *&next) {
69 StringPrinter::StringPrinterBufferPointer<> retval = {nullptr};
70
71 switch (*buffer) {
72 case 0:
73 retval = {"\\0", 2};
74 break;
75 case '\a':
76 retval = {"\\a", 2};
77 break;
78 case '\b':
79 retval = {"\\b", 2};
80 break;
81 case '\f':
82 retval = {"\\f", 2};
83 break;
84 case '\n':
85 retval = {"\\n", 2};
86 break;
87 case '\r':
88 retval = {"\\r", 2};
89 break;
90 case '\t':
91 retval = {"\\t", 2};
92 break;
93 case '\v':
94 retval = {"\\v", 2};
95 break;
96 case '\"':
97 retval = {"\\\"", 2};
98 break;
99 case '\\':
100 retval = {"\\\\", 2};
101 break;
102 default:
103 if (isprint(*buffer))
104 retval = {buffer, 1};
105 else {
106 uint8_t *data = new uint8_t[5];
107 sprintf((char *)data, "\\x%02x", *buffer);
108 retval = {data, 4, [](const uint8_t *c) { delete[] c; }};
109 break;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000110 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000111 }
112
113 next = buffer + 1;
114 return retval;
115}
116
117static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1) {
118 return (c0 - 192) * 64 + (c1 - 128);
119}
120static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1,
121 unsigned char c2) {
122 return (c0 - 224) * 4096 + (c1 - 128) * 64 + (c2 - 128);
123}
124static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1,
125 unsigned char c2, unsigned char c3) {
126 return (c0 - 240) * 262144 + (c2 - 128) * 4096 + (c2 - 128) * 64 + (c3 - 128);
127}
128
129template <>
130StringPrinter::StringPrinterBufferPointer<>
131GetPrintableImpl<StringPrinter::StringElementType::UTF8>(uint8_t *buffer,
132 uint8_t *buffer_end,
133 uint8_t *&next) {
134 StringPrinter::StringPrinterBufferPointer<> retval{nullptr};
135
Justin Lebar90910552016-09-30 00:38:45 +0000136 unsigned utf8_encoded_len = llvm::getNumBytesForUTF8(*buffer);
Kate Stoneb9c1b512016-09-06 20:57:50 +0000137
Zachary Turner5a8ad4592016-10-05 17:07:34 +0000138 if (1u + std::distance(buffer, buffer_end) < utf8_encoded_len) {
Kate Stoneb9c1b512016-09-06 20:57:50 +0000139 // I don't have enough bytes - print whatever I have left
140 retval = {buffer, static_cast<size_t>(1 + buffer_end - buffer)};
141 next = buffer_end + 1;
142 return retval;
143 }
144
145 char32_t codepoint = 0;
146 switch (utf8_encoded_len) {
147 case 1:
148 // this is just an ASCII byte - ask ASCII
149 return GetPrintableImpl<StringPrinter::StringElementType::ASCII>(
150 buffer, buffer_end, next);
151 case 2:
152 codepoint = ConvertUTF8ToCodePoint((unsigned char)*buffer,
153 (unsigned char)*(buffer + 1));
154 break;
155 case 3:
156 codepoint = ConvertUTF8ToCodePoint((unsigned char)*buffer,
157 (unsigned char)*(buffer + 1),
158 (unsigned char)*(buffer + 2));
159 break;
160 case 4:
161 codepoint = ConvertUTF8ToCodePoint(
162 (unsigned char)*buffer, (unsigned char)*(buffer + 1),
163 (unsigned char)*(buffer + 2), (unsigned char)*(buffer + 3));
164 break;
165 default:
166 // this is probably some bogus non-character thing
167 // just print it as-is and hope to sync up again soon
168 retval = {buffer, 1};
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000169 next = buffer + 1;
170 return retval;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000171 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000172
Kate Stoneb9c1b512016-09-06 20:57:50 +0000173 if (codepoint) {
174 switch (codepoint) {
175 case 0:
176 retval = {"\\0", 2};
177 break;
178 case '\a':
179 retval = {"\\a", 2};
180 break;
181 case '\b':
182 retval = {"\\b", 2};
183 break;
184 case '\f':
185 retval = {"\\f", 2};
186 break;
187 case '\n':
188 retval = {"\\n", 2};
189 break;
190 case '\r':
191 retval = {"\\r", 2};
192 break;
193 case '\t':
194 retval = {"\\t", 2};
195 break;
196 case '\v':
197 retval = {"\\v", 2};
198 break;
199 case '\"':
200 retval = {"\\\"", 2};
201 break;
202 case '\\':
203 retval = {"\\\\", 2};
204 break;
205 default:
206 if (isprint(codepoint))
207 retval = {buffer, utf8_encoded_len};
208 else {
209 uint8_t *data = new uint8_t[11];
210 sprintf((char *)data, "\\U%08x", (unsigned)codepoint);
211 retval = {data, 10, [](const uint8_t *c) { delete[] c; }};
212 break;
213 }
214 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000215
Kate Stoneb9c1b512016-09-06 20:57:50 +0000216 next = buffer + utf8_encoded_len;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000217 return retval;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000218 }
219
220 // this should not happen - but just in case.. try to resync at some point
221 retval = {buffer, 1};
222 next = buffer + 1;
223 return retval;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000224}
225
226// Given a sequence of bytes, this function returns:
227// a sequence of bytes to actually print out + a length
228// the following unscanned position of the buffer is in next
Enrico Granataad650a12015-09-09 20:59:49 +0000229static StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000230GetPrintable(StringPrinter::StringElementType type, uint8_t *buffer,
231 uint8_t *buffer_end, uint8_t *&next) {
232 if (!buffer)
233 return {nullptr};
234
235 switch (type) {
236 case StringPrinter::StringElementType::ASCII:
237 return GetPrintableImpl<StringPrinter::StringElementType::ASCII>(
238 buffer, buffer_end, next);
239 case StringPrinter::StringElementType::UTF8:
240 return GetPrintableImpl<StringPrinter::StringElementType::UTF8>(
241 buffer, buffer_end, next);
242 default:
243 return {nullptr};
244 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000245}
246
Enrico Granataac494532015-09-09 22:30:24 +0000247StringPrinter::EscapingHelper
Kate Stoneb9c1b512016-09-06 20:57:50 +0000248StringPrinter::GetDefaultEscapingHelper(GetPrintableElementType elem_type) {
249 switch (elem_type) {
250 case GetPrintableElementType::UTF8:
251 return [](uint8_t *buffer, uint8_t *buffer_end,
252 uint8_t *&next) -> StringPrinter::StringPrinterBufferPointer<> {
253 return GetPrintable(StringPrinter::StringElementType::UTF8, buffer,
254 buffer_end, next);
255 };
256 case GetPrintableElementType::ASCII:
257 return [](uint8_t *buffer, uint8_t *buffer_end,
258 uint8_t *&next) -> StringPrinter::StringPrinterBufferPointer<> {
259 return GetPrintable(StringPrinter::StringElementType::ASCII, buffer,
260 buffer_end, next);
261 };
262 }
263 llvm_unreachable("bad element type");
Enrico Granataac494532015-09-09 22:30:24 +0000264}
265
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000266// use this call if you already have an LLDB-side buffer for the data
Kate Stoneb9c1b512016-09-06 20:57:50 +0000267template <typename SourceDataType>
268static bool DumpUTFBufferToStream(
Justin Lebar90910552016-09-30 00:38:45 +0000269 llvm::ConversionResult (*ConvertFunction)(const SourceDataType **,
270 const SourceDataType *,
271 llvm::UTF8 **, llvm::UTF8 *,
272 llvm::ConversionFlags),
Kate Stoneb9c1b512016-09-06 20:57:50 +0000273 const StringPrinter::ReadBufferAndDumpToStreamOptions &dump_options) {
274 Stream &stream(*dump_options.GetStream());
275 if (dump_options.GetPrefixToken() != 0)
276 stream.Printf("%s", dump_options.GetPrefixToken());
277 if (dump_options.GetQuote() != 0)
278 stream.Printf("%c", dump_options.GetQuote());
279 auto data(dump_options.GetData());
280 auto source_size(dump_options.GetSourceSize());
281 if (data.GetByteSize() && data.GetDataStart() && data.GetDataEnd()) {
282 const int bufferSPSize = data.GetByteSize();
283 if (dump_options.GetSourceSize() == 0) {
284 const int origin_encoding = 8 * sizeof(SourceDataType);
285 source_size = bufferSPSize / (origin_encoding / 4);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000286 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000287
Kate Stoneb9c1b512016-09-06 20:57:50 +0000288 const SourceDataType *data_ptr =
289 (const SourceDataType *)data.GetDataStart();
290 const SourceDataType *data_end_ptr = data_ptr + source_size;
Enrico Granataebdc1ac2014-11-05 21:20:48 +0000291
Kate Stoneb9c1b512016-09-06 20:57:50 +0000292 const bool zero_is_terminator = dump_options.GetBinaryZeroIsTerminator();
Enrico Granataebdc1ac2014-11-05 21:20:48 +0000293
Kate Stoneb9c1b512016-09-06 20:57:50 +0000294 if (zero_is_terminator) {
295 while (data_ptr < data_end_ptr) {
296 if (!*data_ptr) {
297 data_end_ptr = data_ptr;
298 break;
Enrico Granatab7662922015-11-04 00:02:08 +0000299 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000300 data_ptr++;
301 }
302
303 data_ptr = (const SourceDataType *)data.GetDataStart();
Enrico Granatab7662922015-11-04 00:02:08 +0000304 }
Shawn Bestfd137432014-11-04 22:43:34 +0000305
Kate Stoneb9c1b512016-09-06 20:57:50 +0000306 lldb::DataBufferSP utf8_data_buffer_sp;
Justin Lebar90910552016-09-30 00:38:45 +0000307 llvm::UTF8 *utf8_data_ptr = nullptr;
308 llvm::UTF8 *utf8_data_end_ptr = nullptr;
Shawn Bestfd137432014-11-04 22:43:34 +0000309
Kate Stoneb9c1b512016-09-06 20:57:50 +0000310 if (ConvertFunction) {
311 utf8_data_buffer_sp.reset(new DataBufferHeap(4 * bufferSPSize, 0));
Justin Lebar90910552016-09-30 00:38:45 +0000312 utf8_data_ptr = (llvm::UTF8 *)utf8_data_buffer_sp->GetBytes();
Kate Stoneb9c1b512016-09-06 20:57:50 +0000313 utf8_data_end_ptr = utf8_data_ptr + utf8_data_buffer_sp->GetByteSize();
314 ConvertFunction(&data_ptr, data_end_ptr, &utf8_data_ptr,
Justin Lebar90910552016-09-30 00:38:45 +0000315 utf8_data_end_ptr, llvm::lenientConversion);
Kate Stoneb9c1b512016-09-06 20:57:50 +0000316 if (false == zero_is_terminator)
317 utf8_data_end_ptr = utf8_data_ptr;
Justin Lebar90910552016-09-30 00:38:45 +0000318 // needed because the ConvertFunction will change the value of the
319 // data_ptr.
Kate Stoneb9c1b512016-09-06 20:57:50 +0000320 utf8_data_ptr =
Justin Lebar90910552016-09-30 00:38:45 +0000321 (llvm::UTF8 *)utf8_data_buffer_sp->GetBytes();
Kate Stoneb9c1b512016-09-06 20:57:50 +0000322 } else {
323 // just copy the pointers - the cast is necessary to make the compiler
324 // happy
325 // but this should only happen if we are reading UTF8 data
Justin Lebar90910552016-09-30 00:38:45 +0000326 utf8_data_ptr = const_cast<llvm::UTF8 *>(
327 reinterpret_cast<const llvm::UTF8 *>(data_ptr));
328 utf8_data_end_ptr = const_cast<llvm::UTF8 *>(
329 reinterpret_cast<const llvm::UTF8 *>(data_end_ptr));
Kate Stoneb9c1b512016-09-06 20:57:50 +0000330 }
Shawn Bestfd137432014-11-04 22:43:34 +0000331
Kate Stoneb9c1b512016-09-06 20:57:50 +0000332 const bool escape_non_printables = dump_options.GetEscapeNonPrintables();
Enrico Granataac494532015-09-09 22:30:24 +0000333 lldb_private::formatters::StringPrinter::EscapingHelper escaping_callback;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000334 if (escape_non_printables) {
335 if (Language *language = Language::FindPlugin(dump_options.GetLanguage()))
336 escaping_callback = language->GetStringPrinterEscapingHelper(
337 lldb_private::formatters::StringPrinter::GetPrintableElementType::
338 UTF8);
339 else
340 escaping_callback =
341 lldb_private::formatters::StringPrinter::GetDefaultEscapingHelper(
342 lldb_private::formatters::StringPrinter::
343 GetPrintableElementType::UTF8);
Enrico Granataac494532015-09-09 22:30:24 +0000344 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000345
Shawn Bestfd137432014-11-04 22:43:34 +0000346 // since we tend to accept partial data (and even partially malformed data)
347 // we might end up with no NULL terminator before the end_ptr
348 // hence we need to take a slower route and ensure we stay within boundaries
Kate Stoneb9c1b512016-09-06 20:57:50 +0000349 for (; utf8_data_ptr < utf8_data_end_ptr;) {
350 if (zero_is_terminator && !*utf8_data_ptr)
351 break;
Shawn Bestfd137432014-11-04 22:43:34 +0000352
Kate Stoneb9c1b512016-09-06 20:57:50 +0000353 if (escape_non_printables) {
354 uint8_t *next_data = nullptr;
355 auto printable =
356 escaping_callback(utf8_data_ptr, utf8_data_end_ptr, next_data);
357 auto printable_bytes = printable.GetBytes();
358 auto printable_size = printable.GetSize();
359 if (!printable_bytes || !next_data) {
360 // GetPrintable() failed on us - print one byte in a desperate resync
361 // attempt
362 printable_bytes = utf8_data_ptr;
363 printable_size = 1;
364 next_data = utf8_data_ptr + 1;
365 }
366 for (unsigned c = 0; c < printable_size; c++)
367 stream.Printf("%c", *(printable_bytes + c));
368 utf8_data_ptr = (uint8_t *)next_data;
369 } else {
370 stream.Printf("%c", *utf8_data_ptr);
371 utf8_data_ptr++;
372 }
373 }
374 }
375 if (dump_options.GetQuote() != 0)
376 stream.Printf("%c", dump_options.GetQuote());
377 if (dump_options.GetSuffixToken() != 0)
378 stream.Printf("%s", dump_options.GetSuffixToken());
379 if (dump_options.GetIsTruncated())
380 stream.Printf("...");
381 return true;
Shawn Bestfd137432014-11-04 22:43:34 +0000382}
383
Kate Stoneb9c1b512016-09-06 20:57:50 +0000384lldb_private::formatters::StringPrinter::ReadStringAndDumpToStreamOptions::
385 ReadStringAndDumpToStreamOptions(ValueObject &valobj)
386 : ReadStringAndDumpToStreamOptions() {
387 SetEscapeNonPrintables(
388 valobj.GetTargetSP()->GetDebugger().GetEscapeNonPrintables());
389}
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000390
Kate Stoneb9c1b512016-09-06 20:57:50 +0000391lldb_private::formatters::StringPrinter::ReadBufferAndDumpToStreamOptions::
392 ReadBufferAndDumpToStreamOptions(ValueObject &valobj)
393 : ReadBufferAndDumpToStreamOptions() {
394 SetEscapeNonPrintables(
395 valobj.GetTargetSP()->GetDebugger().GetEscapeNonPrintables());
396}
Shawn Bestfd137432014-11-04 22:43:34 +0000397
Kate Stoneb9c1b512016-09-06 20:57:50 +0000398lldb_private::formatters::StringPrinter::ReadBufferAndDumpToStreamOptions::
399 ReadBufferAndDumpToStreamOptions(
400 const ReadStringAndDumpToStreamOptions &options)
401 : ReadBufferAndDumpToStreamOptions() {
402 SetStream(options.GetStream());
403 SetPrefixToken(options.GetPrefixToken());
404 SetSuffixToken(options.GetSuffixToken());
405 SetQuote(options.GetQuote());
406 SetEscapeNonPrintables(options.GetEscapeNonPrintables());
407 SetBinaryZeroIsTerminator(options.GetBinaryZeroIsTerminator());
408 SetLanguage(options.GetLanguage());
409}
Shawn Bestfd137432014-11-04 22:43:34 +0000410
Kate Stoneb9c1b512016-09-06 20:57:50 +0000411namespace lldb_private {
Shawn Bestfd137432014-11-04 22:43:34 +0000412
Kate Stoneb9c1b512016-09-06 20:57:50 +0000413namespace formatters {
Shawn Bestfd137432014-11-04 22:43:34 +0000414
Kate Stoneb9c1b512016-09-06 20:57:50 +0000415template <>
416bool StringPrinter::ReadStringAndDumpToStream<
417 StringPrinter::StringElementType::ASCII>(
418 const ReadStringAndDumpToStreamOptions &options) {
419 assert(options.GetStream() && "need a Stream to print the string to");
420 Error my_error;
Shawn Bestfd137432014-11-04 22:43:34 +0000421
Kate Stoneb9c1b512016-09-06 20:57:50 +0000422 ProcessSP process_sp(options.GetProcessSP());
Shawn Bestfd137432014-11-04 22:43:34 +0000423
Kate Stoneb9c1b512016-09-06 20:57:50 +0000424 if (process_sp.get() == nullptr || options.GetLocation() == 0)
425 return false;
426
427 size_t size;
428 const auto max_size = process_sp->GetTarget().GetMaximumSizeOfStringSummary();
429 bool is_truncated = false;
430
431 if (options.GetSourceSize() == 0)
432 size = max_size;
433 else if (!options.GetIgnoreMaxLength()) {
434 size = options.GetSourceSize();
435 if (size > max_size) {
436 size = max_size;
437 is_truncated = true;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000438 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000439 } else
440 size = options.GetSourceSize();
Shawn Bestfd137432014-11-04 22:43:34 +0000441
Kate Stoneb9c1b512016-09-06 20:57:50 +0000442 lldb::DataBufferSP buffer_sp(new DataBufferHeap(size, 0));
Shawn Bestfd137432014-11-04 22:43:34 +0000443
Kate Stoneb9c1b512016-09-06 20:57:50 +0000444 process_sp->ReadCStringFromMemory(
445 options.GetLocation(), (char *)buffer_sp->GetBytes(), size, my_error);
Shawn Bestfd137432014-11-04 22:43:34 +0000446
Kate Stoneb9c1b512016-09-06 20:57:50 +0000447 if (my_error.Fail())
448 return false;
Shawn Bestfd137432014-11-04 22:43:34 +0000449
Kate Stoneb9c1b512016-09-06 20:57:50 +0000450 const char *prefix_token = options.GetPrefixToken();
451 char quote = options.GetQuote();
Shawn Bestfd137432014-11-04 22:43:34 +0000452
Kate Stoneb9c1b512016-09-06 20:57:50 +0000453 if (prefix_token != 0)
454 options.GetStream()->Printf("%s%c", prefix_token, quote);
455 else if (quote != 0)
456 options.GetStream()->Printf("%c", quote);
457
458 uint8_t *data_end = buffer_sp->GetBytes() + buffer_sp->GetByteSize();
459
460 const bool escape_non_printables = options.GetEscapeNonPrintables();
461 lldb_private::formatters::StringPrinter::EscapingHelper escaping_callback;
462 if (escape_non_printables) {
463 if (Language *language = Language::FindPlugin(options.GetLanguage()))
464 escaping_callback = language->GetStringPrinterEscapingHelper(
465 lldb_private::formatters::StringPrinter::GetPrintableElementType::
466 ASCII);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000467 else
Kate Stoneb9c1b512016-09-06 20:57:50 +0000468 escaping_callback =
469 lldb_private::formatters::StringPrinter::GetDefaultEscapingHelper(
470 lldb_private::formatters::StringPrinter::GetPrintableElementType::
471 ASCII);
472 }
Shawn Bestfd137432014-11-04 22:43:34 +0000473
Kate Stoneb9c1b512016-09-06 20:57:50 +0000474 // since we tend to accept partial data (and even partially malformed data)
475 // we might end up with no NULL terminator before the end_ptr
476 // hence we need to take a slower route and ensure we stay within boundaries
477 for (uint8_t *data = buffer_sp->GetBytes(); *data && (data < data_end);) {
478 if (escape_non_printables) {
479 uint8_t *next_data = nullptr;
480 auto printable = escaping_callback(data, data_end, next_data);
481 auto printable_bytes = printable.GetBytes();
482 auto printable_size = printable.GetSize();
483 if (!printable_bytes || !next_data) {
484 // GetPrintable() failed on us - print one byte in a desperate resync
485 // attempt
486 printable_bytes = data;
487 printable_size = 1;
488 next_data = data + 1;
489 }
490 for (unsigned c = 0; c < printable_size; c++)
491 options.GetStream()->Printf("%c", *(printable_bytes + c));
492 data = (uint8_t *)next_data;
493 } else {
494 options.GetStream()->Printf("%c", *data);
495 data++;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000496 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000497 }
Shawn Bestfd137432014-11-04 22:43:34 +0000498
Kate Stoneb9c1b512016-09-06 20:57:50 +0000499 const char *suffix_token = options.GetSuffixToken();
Shawn Bestfd137432014-11-04 22:43:34 +0000500
Kate Stoneb9c1b512016-09-06 20:57:50 +0000501 if (suffix_token != 0)
502 options.GetStream()->Printf("%c%s", quote, suffix_token);
503 else if (quote != 0)
504 options.GetStream()->Printf("%c", quote);
505
506 if (is_truncated)
507 options.GetStream()->Printf("...");
508
509 return true;
510}
511
512template <typename SourceDataType>
513static bool ReadUTFBufferAndDumpToStream(
514 const StringPrinter::ReadStringAndDumpToStreamOptions &options,
Justin Lebar90910552016-09-30 00:38:45 +0000515 llvm::ConversionResult (*ConvertFunction)(const SourceDataType **,
516 const SourceDataType *,
517 llvm::UTF8 **, llvm::UTF8 *,
518 llvm::ConversionFlags)) {
Kate Stoneb9c1b512016-09-06 20:57:50 +0000519 assert(options.GetStream() && "need a Stream to print the string to");
520
521 if (options.GetLocation() == 0 ||
522 options.GetLocation() == LLDB_INVALID_ADDRESS)
523 return false;
524
525 lldb::ProcessSP process_sp(options.GetProcessSP());
526
527 if (!process_sp)
528 return false;
529
530 const int type_width = sizeof(SourceDataType);
531 const int origin_encoding = 8 * type_width;
532 if (origin_encoding != 8 && origin_encoding != 16 && origin_encoding != 32)
533 return false;
534 // if not UTF8, I need a conversion function to return proper UTF8
535 if (origin_encoding != 8 && !ConvertFunction)
536 return false;
537
538 if (!options.GetStream())
539 return false;
540
541 uint32_t sourceSize = options.GetSourceSize();
542 bool needs_zero_terminator = options.GetNeedsZeroTermination();
543
544 bool is_truncated = false;
545 const auto max_size = process_sp->GetTarget().GetMaximumSizeOfStringSummary();
546
547 if (!sourceSize) {
548 sourceSize = max_size;
549 needs_zero_terminator = true;
550 } else if (!options.GetIgnoreMaxLength()) {
551 if (sourceSize > max_size) {
552 sourceSize = max_size;
553 is_truncated = true;
554 }
555 }
556
557 const int bufferSPSize = sourceSize * type_width;
558
559 lldb::DataBufferSP buffer_sp(new DataBufferHeap(bufferSPSize, 0));
560
561 if (!buffer_sp->GetBytes())
562 return false;
563
564 Error error;
565 char *buffer = reinterpret_cast<char *>(buffer_sp->GetBytes());
566
567 if (needs_zero_terminator)
568 process_sp->ReadStringFromMemory(options.GetLocation(), buffer,
569 bufferSPSize, error, type_width);
570 else
571 process_sp->ReadMemoryFromInferior(options.GetLocation(),
572 (char *)buffer_sp->GetBytes(),
573 bufferSPSize, error);
574
575 if (error.Fail()) {
576 options.GetStream()->Printf("unable to read data");
577 return true;
578 }
579
580 DataExtractor data(buffer_sp, process_sp->GetByteOrder(),
581 process_sp->GetAddressByteSize());
582
583 StringPrinter::ReadBufferAndDumpToStreamOptions dump_options(options);
584 dump_options.SetData(data);
585 dump_options.SetSourceSize(sourceSize);
586 dump_options.SetIsTruncated(is_truncated);
587
588 return DumpUTFBufferToStream(ConvertFunction, dump_options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000589}
590
591template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000592bool StringPrinter::ReadStringAndDumpToStream<
593 StringPrinter::StringElementType::UTF8>(
594 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000595 return ReadUTFBufferAndDumpToStream<llvm::UTF8>(options, nullptr);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000596}
597
598template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000599bool StringPrinter::ReadStringAndDumpToStream<
600 StringPrinter::StringElementType::UTF16>(
601 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000602 return ReadUTFBufferAndDumpToStream<llvm::UTF16>(options,
603 llvm::ConvertUTF16toUTF8);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000604}
605
606template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000607bool StringPrinter::ReadStringAndDumpToStream<
608 StringPrinter::StringElementType::UTF32>(
609 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000610 return ReadUTFBufferAndDumpToStream<llvm::UTF32>(options,
611 llvm::ConvertUTF32toUTF8);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000612}
613
614template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000615bool StringPrinter::ReadBufferAndDumpToStream<
616 StringPrinter::StringElementType::UTF8>(
617 const ReadBufferAndDumpToStreamOptions &options) {
618 assert(options.GetStream() && "need a Stream to print the string to");
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000619
Justin Lebar90910552016-09-30 00:38:45 +0000620 return DumpUTFBufferToStream<llvm::UTF8>(nullptr, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000621}
622
623template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000624bool StringPrinter::ReadBufferAndDumpToStream<
625 StringPrinter::StringElementType::ASCII>(
626 const ReadBufferAndDumpToStreamOptions &options) {
627 // treat ASCII the same as UTF8
628 // FIXME: can we optimize ASCII some more?
629 return ReadBufferAndDumpToStream<StringElementType::UTF8>(options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000630}
631
632template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000633bool StringPrinter::ReadBufferAndDumpToStream<
634 StringPrinter::StringElementType::UTF16>(
635 const ReadBufferAndDumpToStreamOptions &options) {
636 assert(options.GetStream() && "need a Stream to print the string to");
Shawn Bestfd137432014-11-04 22:43:34 +0000637
Justin Lebar90910552016-09-30 00:38:45 +0000638 return DumpUTFBufferToStream(llvm::ConvertUTF16toUTF8, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000639}
640
641template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000642bool StringPrinter::ReadBufferAndDumpToStream<
643 StringPrinter::StringElementType::UTF32>(
644 const ReadBufferAndDumpToStreamOptions &options) {
645 assert(options.GetStream() && "need a Stream to print the string to");
Shawn Bestfd137432014-11-04 22:43:34 +0000646
Justin Lebar90910552016-09-30 00:38:45 +0000647 return DumpUTFBufferToStream(llvm::ConvertUTF32toUTF8, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000648}
Shawn Bestfd137432014-11-04 22:43:34 +0000649
650} // namespace formatters
651
652} // namespace lldb_private