blob: 6a000f0f6431481b628a3a82951a266983a9b204 [file] [log] [blame]
Kate Stoneb9c1b512016-09-06 20:57:50 +00001//===-- StringPrinter.cpp ----------------------------------------*- C++
2//-*-===//
Enrico Granataca6c8ee2014-10-30 01:45:39 +00003//
4// The LLVM Compiler Infrastructure
5//
6// This file is distributed under the University of Illinois Open Source
7// License. See LICENSE.TXT for details.
8//
9//===----------------------------------------------------------------------===//
10
11#include "lldb/DataFormatters/StringPrinter.h"
12
Enrico Granataebdc1ac2014-11-05 21:20:48 +000013#include "lldb/Core/Debugger.h"
Enrico Granataebdc1ac2014-11-05 21:20:48 +000014#include "lldb/Core/ValueObject.h"
Enrico Granataac494532015-09-09 22:30:24 +000015#include "lldb/Target/Language.h"
Enrico Granataca6c8ee2014-10-30 01:45:39 +000016#include "lldb/Target/Process.h"
17#include "lldb/Target/Target.h"
Zachary Turner97206d52017-05-12 04:51:55 +000018#include "lldb/Utility/Status.h"
Enrico Granataca6c8ee2014-10-30 01:45:39 +000019
20#include "llvm/Support/ConvertUTF.h"
21
Enrico Granataca6c8ee2014-10-30 01:45:39 +000022#include <ctype.h>
Enrico Granataca6c8ee2014-10-30 01:45:39 +000023#include <locale>
24
25using namespace lldb;
26using namespace lldb_private;
27using namespace lldb_private::formatters;
28
Adrian Prantl05097242018-04-30 16:49:04 +000029// we define this for all values of type but only implement it for those we
30// care about that's good because we get linker errors for any unsupported type
Enrico Granataac494532015-09-09 22:30:24 +000031template <lldb_private::formatters::StringPrinter::StringElementType type>
Enrico Granataad650a12015-09-09 20:59:49 +000032static StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +000033GetPrintableImpl(uint8_t *buffer, uint8_t *buffer_end, uint8_t *&next);
Enrico Granataca6c8ee2014-10-30 01:45:39 +000034
35// mimic isprint() for Unicode codepoints
Kate Stoneb9c1b512016-09-06 20:57:50 +000036static bool isprint(char32_t codepoint) {
37 if (codepoint <= 0x1F || codepoint == 0x7F) // C0
38 {
39 return false;
40 }
41 if (codepoint >= 0x80 && codepoint <= 0x9F) // C1
42 {
43 return false;
44 }
45 if (codepoint == 0x2028 || codepoint == 0x2029) // line/paragraph separators
46 {
47 return false;
48 }
49 if (codepoint == 0x200E || codepoint == 0x200F ||
50 (codepoint >= 0x202A &&
51 codepoint <= 0x202E)) // bidirectional text control
52 {
53 return false;
54 }
55 if (codepoint >= 0xFFF9 &&
56 codepoint <= 0xFFFF) // interlinears and generally specials
57 {
58 return false;
59 }
60 return true;
Enrico Granataca6c8ee2014-10-30 01:45:39 +000061}
62
63template <>
Enrico Granataad650a12015-09-09 20:59:49 +000064StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +000065GetPrintableImpl<StringPrinter::StringElementType::ASCII>(uint8_t *buffer,
66 uint8_t *buffer_end,
67 uint8_t *&next) {
68 StringPrinter::StringPrinterBufferPointer<> retval = {nullptr};
69
70 switch (*buffer) {
71 case 0:
72 retval = {"\\0", 2};
73 break;
74 case '\a':
75 retval = {"\\a", 2};
76 break;
77 case '\b':
78 retval = {"\\b", 2};
79 break;
80 case '\f':
81 retval = {"\\f", 2};
82 break;
83 case '\n':
84 retval = {"\\n", 2};
85 break;
86 case '\r':
87 retval = {"\\r", 2};
88 break;
89 case '\t':
90 retval = {"\\t", 2};
91 break;
92 case '\v':
93 retval = {"\\v", 2};
94 break;
95 case '\"':
96 retval = {"\\\"", 2};
97 break;
98 case '\\':
99 retval = {"\\\\", 2};
100 break;
101 default:
102 if (isprint(*buffer))
103 retval = {buffer, 1};
104 else {
105 uint8_t *data = new uint8_t[5];
106 sprintf((char *)data, "\\x%02x", *buffer);
107 retval = {data, 4, [](const uint8_t *c) { delete[] c; }};
108 break;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000109 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000110 }
111
112 next = buffer + 1;
113 return retval;
114}
115
116static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1) {
117 return (c0 - 192) * 64 + (c1 - 128);
118}
119static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1,
120 unsigned char c2) {
121 return (c0 - 224) * 4096 + (c1 - 128) * 64 + (c2 - 128);
122}
123static char32_t ConvertUTF8ToCodePoint(unsigned char c0, unsigned char c1,
124 unsigned char c2, unsigned char c3) {
125 return (c0 - 240) * 262144 + (c2 - 128) * 4096 + (c2 - 128) * 64 + (c3 - 128);
126}
127
128template <>
129StringPrinter::StringPrinterBufferPointer<>
130GetPrintableImpl<StringPrinter::StringElementType::UTF8>(uint8_t *buffer,
131 uint8_t *buffer_end,
132 uint8_t *&next) {
133 StringPrinter::StringPrinterBufferPointer<> retval{nullptr};
134
Justin Lebar90910552016-09-30 00:38:45 +0000135 unsigned utf8_encoded_len = llvm::getNumBytesForUTF8(*buffer);
Kate Stoneb9c1b512016-09-06 20:57:50 +0000136
Zachary Turner5a8ad4592016-10-05 17:07:34 +0000137 if (1u + std::distance(buffer, buffer_end) < utf8_encoded_len) {
Kate Stoneb9c1b512016-09-06 20:57:50 +0000138 // I don't have enough bytes - print whatever I have left
139 retval = {buffer, static_cast<size_t>(1 + buffer_end - buffer)};
140 next = buffer_end + 1;
141 return retval;
142 }
143
144 char32_t codepoint = 0;
145 switch (utf8_encoded_len) {
146 case 1:
147 // this is just an ASCII byte - ask ASCII
148 return GetPrintableImpl<StringPrinter::StringElementType::ASCII>(
149 buffer, buffer_end, next);
150 case 2:
151 codepoint = ConvertUTF8ToCodePoint((unsigned char)*buffer,
152 (unsigned char)*(buffer + 1));
153 break;
154 case 3:
155 codepoint = ConvertUTF8ToCodePoint((unsigned char)*buffer,
156 (unsigned char)*(buffer + 1),
157 (unsigned char)*(buffer + 2));
158 break;
159 case 4:
160 codepoint = ConvertUTF8ToCodePoint(
161 (unsigned char)*buffer, (unsigned char)*(buffer + 1),
162 (unsigned char)*(buffer + 2), (unsigned char)*(buffer + 3));
163 break;
164 default:
Adrian Prantl05097242018-04-30 16:49:04 +0000165 // this is probably some bogus non-character thing just print it as-is and
166 // hope to sync up again soon
Kate Stoneb9c1b512016-09-06 20:57:50 +0000167 retval = {buffer, 1};
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000168 next = buffer + 1;
169 return retval;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000170 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000171
Kate Stoneb9c1b512016-09-06 20:57:50 +0000172 if (codepoint) {
173 switch (codepoint) {
174 case 0:
175 retval = {"\\0", 2};
176 break;
177 case '\a':
178 retval = {"\\a", 2};
179 break;
180 case '\b':
181 retval = {"\\b", 2};
182 break;
183 case '\f':
184 retval = {"\\f", 2};
185 break;
186 case '\n':
187 retval = {"\\n", 2};
188 break;
189 case '\r':
190 retval = {"\\r", 2};
191 break;
192 case '\t':
193 retval = {"\\t", 2};
194 break;
195 case '\v':
196 retval = {"\\v", 2};
197 break;
198 case '\"':
199 retval = {"\\\"", 2};
200 break;
201 case '\\':
202 retval = {"\\\\", 2};
203 break;
204 default:
205 if (isprint(codepoint))
206 retval = {buffer, utf8_encoded_len};
207 else {
208 uint8_t *data = new uint8_t[11];
209 sprintf((char *)data, "\\U%08x", (unsigned)codepoint);
210 retval = {data, 10, [](const uint8_t *c) { delete[] c; }};
211 break;
212 }
213 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000214
Kate Stoneb9c1b512016-09-06 20:57:50 +0000215 next = buffer + utf8_encoded_len;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000216 return retval;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000217 }
218
219 // this should not happen - but just in case.. try to resync at some point
220 retval = {buffer, 1};
221 next = buffer + 1;
222 return retval;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000223}
224
Adrian Prantl05097242018-04-30 16:49:04 +0000225// Given a sequence of bytes, this function returns: a sequence of bytes to
226// actually print out + a length the following unscanned position of the buffer
227// is in next
Enrico Granataad650a12015-09-09 20:59:49 +0000228static StringPrinter::StringPrinterBufferPointer<>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000229GetPrintable(StringPrinter::StringElementType type, uint8_t *buffer,
230 uint8_t *buffer_end, uint8_t *&next) {
231 if (!buffer)
232 return {nullptr};
233
234 switch (type) {
235 case StringPrinter::StringElementType::ASCII:
236 return GetPrintableImpl<StringPrinter::StringElementType::ASCII>(
237 buffer, buffer_end, next);
238 case StringPrinter::StringElementType::UTF8:
239 return GetPrintableImpl<StringPrinter::StringElementType::UTF8>(
240 buffer, buffer_end, next);
241 default:
242 return {nullptr};
243 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000244}
245
Enrico Granataac494532015-09-09 22:30:24 +0000246StringPrinter::EscapingHelper
Kate Stoneb9c1b512016-09-06 20:57:50 +0000247StringPrinter::GetDefaultEscapingHelper(GetPrintableElementType elem_type) {
248 switch (elem_type) {
249 case GetPrintableElementType::UTF8:
250 return [](uint8_t *buffer, uint8_t *buffer_end,
251 uint8_t *&next) -> StringPrinter::StringPrinterBufferPointer<> {
252 return GetPrintable(StringPrinter::StringElementType::UTF8, buffer,
253 buffer_end, next);
254 };
255 case GetPrintableElementType::ASCII:
256 return [](uint8_t *buffer, uint8_t *buffer_end,
257 uint8_t *&next) -> StringPrinter::StringPrinterBufferPointer<> {
258 return GetPrintable(StringPrinter::StringElementType::ASCII, buffer,
259 buffer_end, next);
260 };
261 }
262 llvm_unreachable("bad element type");
Enrico Granataac494532015-09-09 22:30:24 +0000263}
264
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000265// use this call if you already have an LLDB-side buffer for the data
Kate Stoneb9c1b512016-09-06 20:57:50 +0000266template <typename SourceDataType>
267static bool DumpUTFBufferToStream(
Justin Lebar90910552016-09-30 00:38:45 +0000268 llvm::ConversionResult (*ConvertFunction)(const SourceDataType **,
269 const SourceDataType *,
270 llvm::UTF8 **, llvm::UTF8 *,
271 llvm::ConversionFlags),
Kate Stoneb9c1b512016-09-06 20:57:50 +0000272 const StringPrinter::ReadBufferAndDumpToStreamOptions &dump_options) {
273 Stream &stream(*dump_options.GetStream());
274 if (dump_options.GetPrefixToken() != 0)
275 stream.Printf("%s", dump_options.GetPrefixToken());
276 if (dump_options.GetQuote() != 0)
277 stream.Printf("%c", dump_options.GetQuote());
278 auto data(dump_options.GetData());
279 auto source_size(dump_options.GetSourceSize());
280 if (data.GetByteSize() && data.GetDataStart() && data.GetDataEnd()) {
281 const int bufferSPSize = data.GetByteSize();
282 if (dump_options.GetSourceSize() == 0) {
283 const int origin_encoding = 8 * sizeof(SourceDataType);
284 source_size = bufferSPSize / (origin_encoding / 4);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000285 }
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000286
Kate Stoneb9c1b512016-09-06 20:57:50 +0000287 const SourceDataType *data_ptr =
288 (const SourceDataType *)data.GetDataStart();
289 const SourceDataType *data_end_ptr = data_ptr + source_size;
Enrico Granataebdc1ac2014-11-05 21:20:48 +0000290
Kate Stoneb9c1b512016-09-06 20:57:50 +0000291 const bool zero_is_terminator = dump_options.GetBinaryZeroIsTerminator();
Enrico Granataebdc1ac2014-11-05 21:20:48 +0000292
Kate Stoneb9c1b512016-09-06 20:57:50 +0000293 if (zero_is_terminator) {
294 while (data_ptr < data_end_ptr) {
295 if (!*data_ptr) {
296 data_end_ptr = data_ptr;
297 break;
Enrico Granatab7662922015-11-04 00:02:08 +0000298 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000299 data_ptr++;
300 }
301
302 data_ptr = (const SourceDataType *)data.GetDataStart();
Enrico Granatab7662922015-11-04 00:02:08 +0000303 }
Shawn Bestfd137432014-11-04 22:43:34 +0000304
Kate Stoneb9c1b512016-09-06 20:57:50 +0000305 lldb::DataBufferSP utf8_data_buffer_sp;
Justin Lebar90910552016-09-30 00:38:45 +0000306 llvm::UTF8 *utf8_data_ptr = nullptr;
307 llvm::UTF8 *utf8_data_end_ptr = nullptr;
Shawn Bestfd137432014-11-04 22:43:34 +0000308
Kate Stoneb9c1b512016-09-06 20:57:50 +0000309 if (ConvertFunction) {
310 utf8_data_buffer_sp.reset(new DataBufferHeap(4 * bufferSPSize, 0));
Justin Lebar90910552016-09-30 00:38:45 +0000311 utf8_data_ptr = (llvm::UTF8 *)utf8_data_buffer_sp->GetBytes();
Kate Stoneb9c1b512016-09-06 20:57:50 +0000312 utf8_data_end_ptr = utf8_data_ptr + utf8_data_buffer_sp->GetByteSize();
313 ConvertFunction(&data_ptr, data_end_ptr, &utf8_data_ptr,
Justin Lebar90910552016-09-30 00:38:45 +0000314 utf8_data_end_ptr, llvm::lenientConversion);
Jonas Devliegherea6682a42018-12-15 00:15:33 +0000315 if (!zero_is_terminator)
Kate Stoneb9c1b512016-09-06 20:57:50 +0000316 utf8_data_end_ptr = utf8_data_ptr;
Justin Lebar90910552016-09-30 00:38:45 +0000317 // needed because the ConvertFunction will change the value of the
318 // data_ptr.
Kate Stoneb9c1b512016-09-06 20:57:50 +0000319 utf8_data_ptr =
Justin Lebar90910552016-09-30 00:38:45 +0000320 (llvm::UTF8 *)utf8_data_buffer_sp->GetBytes();
Kate Stoneb9c1b512016-09-06 20:57:50 +0000321 } else {
322 // just copy the pointers - the cast is necessary to make the compiler
Adrian Prantl05097242018-04-30 16:49:04 +0000323 // happy but this should only happen if we are reading UTF8 data
Justin Lebar90910552016-09-30 00:38:45 +0000324 utf8_data_ptr = const_cast<llvm::UTF8 *>(
325 reinterpret_cast<const llvm::UTF8 *>(data_ptr));
326 utf8_data_end_ptr = const_cast<llvm::UTF8 *>(
327 reinterpret_cast<const llvm::UTF8 *>(data_end_ptr));
Kate Stoneb9c1b512016-09-06 20:57:50 +0000328 }
Shawn Bestfd137432014-11-04 22:43:34 +0000329
Kate Stoneb9c1b512016-09-06 20:57:50 +0000330 const bool escape_non_printables = dump_options.GetEscapeNonPrintables();
Enrico Granataac494532015-09-09 22:30:24 +0000331 lldb_private::formatters::StringPrinter::EscapingHelper escaping_callback;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000332 if (escape_non_printables) {
333 if (Language *language = Language::FindPlugin(dump_options.GetLanguage()))
334 escaping_callback = language->GetStringPrinterEscapingHelper(
335 lldb_private::formatters::StringPrinter::GetPrintableElementType::
336 UTF8);
337 else
338 escaping_callback =
339 lldb_private::formatters::StringPrinter::GetDefaultEscapingHelper(
340 lldb_private::formatters::StringPrinter::
341 GetPrintableElementType::UTF8);
Enrico Granataac494532015-09-09 22:30:24 +0000342 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000343
Shawn Bestfd137432014-11-04 22:43:34 +0000344 // since we tend to accept partial data (and even partially malformed data)
Adrian Prantl05097242018-04-30 16:49:04 +0000345 // we might end up with no NULL terminator before the end_ptr hence we need
346 // to take a slower route and ensure we stay within boundaries
Kate Stoneb9c1b512016-09-06 20:57:50 +0000347 for (; utf8_data_ptr < utf8_data_end_ptr;) {
348 if (zero_is_terminator && !*utf8_data_ptr)
349 break;
Shawn Bestfd137432014-11-04 22:43:34 +0000350
Kate Stoneb9c1b512016-09-06 20:57:50 +0000351 if (escape_non_printables) {
352 uint8_t *next_data = nullptr;
353 auto printable =
354 escaping_callback(utf8_data_ptr, utf8_data_end_ptr, next_data);
355 auto printable_bytes = printable.GetBytes();
356 auto printable_size = printable.GetSize();
357 if (!printable_bytes || !next_data) {
358 // GetPrintable() failed on us - print one byte in a desperate resync
359 // attempt
360 printable_bytes = utf8_data_ptr;
361 printable_size = 1;
362 next_data = utf8_data_ptr + 1;
363 }
364 for (unsigned c = 0; c < printable_size; c++)
365 stream.Printf("%c", *(printable_bytes + c));
366 utf8_data_ptr = (uint8_t *)next_data;
367 } else {
368 stream.Printf("%c", *utf8_data_ptr);
369 utf8_data_ptr++;
370 }
371 }
372 }
373 if (dump_options.GetQuote() != 0)
374 stream.Printf("%c", dump_options.GetQuote());
375 if (dump_options.GetSuffixToken() != 0)
376 stream.Printf("%s", dump_options.GetSuffixToken());
377 if (dump_options.GetIsTruncated())
378 stream.Printf("...");
379 return true;
Shawn Bestfd137432014-11-04 22:43:34 +0000380}
381
Kate Stoneb9c1b512016-09-06 20:57:50 +0000382lldb_private::formatters::StringPrinter::ReadStringAndDumpToStreamOptions::
383 ReadStringAndDumpToStreamOptions(ValueObject &valobj)
384 : ReadStringAndDumpToStreamOptions() {
385 SetEscapeNonPrintables(
386 valobj.GetTargetSP()->GetDebugger().GetEscapeNonPrintables());
387}
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000388
Kate Stoneb9c1b512016-09-06 20:57:50 +0000389lldb_private::formatters::StringPrinter::ReadBufferAndDumpToStreamOptions::
390 ReadBufferAndDumpToStreamOptions(ValueObject &valobj)
391 : ReadBufferAndDumpToStreamOptions() {
392 SetEscapeNonPrintables(
393 valobj.GetTargetSP()->GetDebugger().GetEscapeNonPrintables());
394}
Shawn Bestfd137432014-11-04 22:43:34 +0000395
Kate Stoneb9c1b512016-09-06 20:57:50 +0000396lldb_private::formatters::StringPrinter::ReadBufferAndDumpToStreamOptions::
397 ReadBufferAndDumpToStreamOptions(
398 const ReadStringAndDumpToStreamOptions &options)
399 : ReadBufferAndDumpToStreamOptions() {
400 SetStream(options.GetStream());
401 SetPrefixToken(options.GetPrefixToken());
402 SetSuffixToken(options.GetSuffixToken());
403 SetQuote(options.GetQuote());
404 SetEscapeNonPrintables(options.GetEscapeNonPrintables());
405 SetBinaryZeroIsTerminator(options.GetBinaryZeroIsTerminator());
406 SetLanguage(options.GetLanguage());
407}
Shawn Bestfd137432014-11-04 22:43:34 +0000408
Kate Stoneb9c1b512016-09-06 20:57:50 +0000409namespace lldb_private {
Shawn Bestfd137432014-11-04 22:43:34 +0000410
Kate Stoneb9c1b512016-09-06 20:57:50 +0000411namespace formatters {
Shawn Bestfd137432014-11-04 22:43:34 +0000412
Kate Stoneb9c1b512016-09-06 20:57:50 +0000413template <>
414bool StringPrinter::ReadStringAndDumpToStream<
415 StringPrinter::StringElementType::ASCII>(
416 const ReadStringAndDumpToStreamOptions &options) {
417 assert(options.GetStream() && "need a Stream to print the string to");
Zachary Turner97206d52017-05-12 04:51:55 +0000418 Status my_error;
Shawn Bestfd137432014-11-04 22:43:34 +0000419
Kate Stoneb9c1b512016-09-06 20:57:50 +0000420 ProcessSP process_sp(options.GetProcessSP());
Shawn Bestfd137432014-11-04 22:43:34 +0000421
Kate Stoneb9c1b512016-09-06 20:57:50 +0000422 if (process_sp.get() == nullptr || options.GetLocation() == 0)
423 return false;
424
425 size_t size;
426 const auto max_size = process_sp->GetTarget().GetMaximumSizeOfStringSummary();
427 bool is_truncated = false;
428
429 if (options.GetSourceSize() == 0)
430 size = max_size;
431 else if (!options.GetIgnoreMaxLength()) {
432 size = options.GetSourceSize();
433 if (size > max_size) {
434 size = max_size;
435 is_truncated = true;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000436 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000437 } else
438 size = options.GetSourceSize();
Shawn Bestfd137432014-11-04 22:43:34 +0000439
Kate Stoneb9c1b512016-09-06 20:57:50 +0000440 lldb::DataBufferSP buffer_sp(new DataBufferHeap(size, 0));
Shawn Bestfd137432014-11-04 22:43:34 +0000441
Kate Stoneb9c1b512016-09-06 20:57:50 +0000442 process_sp->ReadCStringFromMemory(
443 options.GetLocation(), (char *)buffer_sp->GetBytes(), size, my_error);
Shawn Bestfd137432014-11-04 22:43:34 +0000444
Kate Stoneb9c1b512016-09-06 20:57:50 +0000445 if (my_error.Fail())
446 return false;
Shawn Bestfd137432014-11-04 22:43:34 +0000447
Kate Stoneb9c1b512016-09-06 20:57:50 +0000448 const char *prefix_token = options.GetPrefixToken();
449 char quote = options.GetQuote();
Shawn Bestfd137432014-11-04 22:43:34 +0000450
Kate Stoneb9c1b512016-09-06 20:57:50 +0000451 if (prefix_token != 0)
452 options.GetStream()->Printf("%s%c", prefix_token, quote);
453 else if (quote != 0)
454 options.GetStream()->Printf("%c", quote);
455
456 uint8_t *data_end = buffer_sp->GetBytes() + buffer_sp->GetByteSize();
457
458 const bool escape_non_printables = options.GetEscapeNonPrintables();
459 lldb_private::formatters::StringPrinter::EscapingHelper escaping_callback;
460 if (escape_non_printables) {
461 if (Language *language = Language::FindPlugin(options.GetLanguage()))
462 escaping_callback = language->GetStringPrinterEscapingHelper(
463 lldb_private::formatters::StringPrinter::GetPrintableElementType::
464 ASCII);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000465 else
Kate Stoneb9c1b512016-09-06 20:57:50 +0000466 escaping_callback =
467 lldb_private::formatters::StringPrinter::GetDefaultEscapingHelper(
468 lldb_private::formatters::StringPrinter::GetPrintableElementType::
469 ASCII);
470 }
Shawn Bestfd137432014-11-04 22:43:34 +0000471
Kate Stoneb9c1b512016-09-06 20:57:50 +0000472 // since we tend to accept partial data (and even partially malformed data)
Adrian Prantl05097242018-04-30 16:49:04 +0000473 // we might end up with no NULL terminator before the end_ptr hence we need
474 // to take a slower route and ensure we stay within boundaries
Kate Stoneb9c1b512016-09-06 20:57:50 +0000475 for (uint8_t *data = buffer_sp->GetBytes(); *data && (data < data_end);) {
476 if (escape_non_printables) {
477 uint8_t *next_data = nullptr;
478 auto printable = escaping_callback(data, data_end, next_data);
479 auto printable_bytes = printable.GetBytes();
480 auto printable_size = printable.GetSize();
481 if (!printable_bytes || !next_data) {
482 // GetPrintable() failed on us - print one byte in a desperate resync
483 // attempt
484 printable_bytes = data;
485 printable_size = 1;
486 next_data = data + 1;
487 }
488 for (unsigned c = 0; c < printable_size; c++)
489 options.GetStream()->Printf("%c", *(printable_bytes + c));
490 data = (uint8_t *)next_data;
491 } else {
492 options.GetStream()->Printf("%c", *data);
493 data++;
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000494 }
Kate Stoneb9c1b512016-09-06 20:57:50 +0000495 }
Shawn Bestfd137432014-11-04 22:43:34 +0000496
Kate Stoneb9c1b512016-09-06 20:57:50 +0000497 const char *suffix_token = options.GetSuffixToken();
Shawn Bestfd137432014-11-04 22:43:34 +0000498
Kate Stoneb9c1b512016-09-06 20:57:50 +0000499 if (suffix_token != 0)
500 options.GetStream()->Printf("%c%s", quote, suffix_token);
501 else if (quote != 0)
502 options.GetStream()->Printf("%c", quote);
503
504 if (is_truncated)
505 options.GetStream()->Printf("...");
506
507 return true;
508}
509
510template <typename SourceDataType>
511static bool ReadUTFBufferAndDumpToStream(
512 const StringPrinter::ReadStringAndDumpToStreamOptions &options,
Justin Lebar90910552016-09-30 00:38:45 +0000513 llvm::ConversionResult (*ConvertFunction)(const SourceDataType **,
514 const SourceDataType *,
515 llvm::UTF8 **, llvm::UTF8 *,
516 llvm::ConversionFlags)) {
Kate Stoneb9c1b512016-09-06 20:57:50 +0000517 assert(options.GetStream() && "need a Stream to print the string to");
518
519 if (options.GetLocation() == 0 ||
520 options.GetLocation() == LLDB_INVALID_ADDRESS)
521 return false;
522
523 lldb::ProcessSP process_sp(options.GetProcessSP());
524
525 if (!process_sp)
526 return false;
527
528 const int type_width = sizeof(SourceDataType);
529 const int origin_encoding = 8 * type_width;
530 if (origin_encoding != 8 && origin_encoding != 16 && origin_encoding != 32)
531 return false;
532 // if not UTF8, I need a conversion function to return proper UTF8
533 if (origin_encoding != 8 && !ConvertFunction)
534 return false;
535
536 if (!options.GetStream())
537 return false;
538
539 uint32_t sourceSize = options.GetSourceSize();
540 bool needs_zero_terminator = options.GetNeedsZeroTermination();
541
542 bool is_truncated = false;
543 const auto max_size = process_sp->GetTarget().GetMaximumSizeOfStringSummary();
544
545 if (!sourceSize) {
546 sourceSize = max_size;
547 needs_zero_terminator = true;
548 } else if (!options.GetIgnoreMaxLength()) {
549 if (sourceSize > max_size) {
550 sourceSize = max_size;
551 is_truncated = true;
552 }
553 }
554
555 const int bufferSPSize = sourceSize * type_width;
556
557 lldb::DataBufferSP buffer_sp(new DataBufferHeap(bufferSPSize, 0));
558
559 if (!buffer_sp->GetBytes())
560 return false;
561
Zachary Turner97206d52017-05-12 04:51:55 +0000562 Status error;
Kate Stoneb9c1b512016-09-06 20:57:50 +0000563 char *buffer = reinterpret_cast<char *>(buffer_sp->GetBytes());
564
565 if (needs_zero_terminator)
566 process_sp->ReadStringFromMemory(options.GetLocation(), buffer,
567 bufferSPSize, error, type_width);
568 else
569 process_sp->ReadMemoryFromInferior(options.GetLocation(),
570 (char *)buffer_sp->GetBytes(),
571 bufferSPSize, error);
572
573 if (error.Fail()) {
574 options.GetStream()->Printf("unable to read data");
575 return true;
576 }
577
578 DataExtractor data(buffer_sp, process_sp->GetByteOrder(),
579 process_sp->GetAddressByteSize());
580
581 StringPrinter::ReadBufferAndDumpToStreamOptions dump_options(options);
582 dump_options.SetData(data);
583 dump_options.SetSourceSize(sourceSize);
584 dump_options.SetIsTruncated(is_truncated);
585
586 return DumpUTFBufferToStream(ConvertFunction, dump_options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000587}
588
589template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000590bool StringPrinter::ReadStringAndDumpToStream<
591 StringPrinter::StringElementType::UTF8>(
592 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000593 return ReadUTFBufferAndDumpToStream<llvm::UTF8>(options, nullptr);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000594}
595
596template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000597bool StringPrinter::ReadStringAndDumpToStream<
598 StringPrinter::StringElementType::UTF16>(
599 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000600 return ReadUTFBufferAndDumpToStream<llvm::UTF16>(options,
601 llvm::ConvertUTF16toUTF8);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000602}
603
604template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000605bool StringPrinter::ReadStringAndDumpToStream<
606 StringPrinter::StringElementType::UTF32>(
607 const ReadStringAndDumpToStreamOptions &options) {
Justin Lebar90910552016-09-30 00:38:45 +0000608 return ReadUTFBufferAndDumpToStream<llvm::UTF32>(options,
609 llvm::ConvertUTF32toUTF8);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000610}
611
612template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000613bool StringPrinter::ReadBufferAndDumpToStream<
614 StringPrinter::StringElementType::UTF8>(
615 const ReadBufferAndDumpToStreamOptions &options) {
616 assert(options.GetStream() && "need a Stream to print the string to");
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000617
Justin Lebar90910552016-09-30 00:38:45 +0000618 return DumpUTFBufferToStream<llvm::UTF8>(nullptr, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000619}
620
621template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000622bool StringPrinter::ReadBufferAndDumpToStream<
623 StringPrinter::StringElementType::ASCII>(
624 const ReadBufferAndDumpToStreamOptions &options) {
625 // treat ASCII the same as UTF8
626 // FIXME: can we optimize ASCII some more?
627 return ReadBufferAndDumpToStream<StringElementType::UTF8>(options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000628}
629
630template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000631bool StringPrinter::ReadBufferAndDumpToStream<
632 StringPrinter::StringElementType::UTF16>(
633 const ReadBufferAndDumpToStreamOptions &options) {
634 assert(options.GetStream() && "need a Stream to print the string to");
Shawn Bestfd137432014-11-04 22:43:34 +0000635
Justin Lebar90910552016-09-30 00:38:45 +0000636 return DumpUTFBufferToStream(llvm::ConvertUTF16toUTF8, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000637}
638
639template <>
Kate Stoneb9c1b512016-09-06 20:57:50 +0000640bool StringPrinter::ReadBufferAndDumpToStream<
641 StringPrinter::StringElementType::UTF32>(
642 const ReadBufferAndDumpToStreamOptions &options) {
643 assert(options.GetStream() && "need a Stream to print the string to");
Shawn Bestfd137432014-11-04 22:43:34 +0000644
Justin Lebar90910552016-09-30 00:38:45 +0000645 return DumpUTFBufferToStream(llvm::ConvertUTF32toUTF8, options);
Enrico Granataca6c8ee2014-10-30 01:45:39 +0000646}
Shawn Bestfd137432014-11-04 22:43:34 +0000647
648} // namespace formatters
649
650} // namespace lldb_private