2021-04-18 17:35:40 -04:00
|
|
|
/*
|
2022-01-31 13:07:22 -05:00
|
|
|
* Copyright (c) 2021, Tim Flynn <trflynn89@serenityos.org>
|
2021-06-21 11:20:09 -04:00
|
|
|
* Copyright (c) 2021, Jan de Visser <jan@de-visser.net>
|
2021-04-18 17:35:40 -04:00
|
|
|
*
|
2021-04-22 01:24:48 -07:00
|
|
|
* SPDX-License-Identifier: BSD-2-Clause
|
2021-04-18 17:35:40 -04:00
|
|
|
*/
|
|
|
|
|
|
|
|
|
|
#pragma once
|
|
|
|
|
|
|
|
|
|
#include "Token.h"
|
2023-12-16 17:49:34 +03:30
|
|
|
#include <AK/ByteString.h>
|
2021-04-18 17:35:40 -04:00
|
|
|
#include <AK/HashMap.h>
|
|
|
|
|
#include <AK/StringView.h>
|
|
|
|
|
|
2021-06-21 10:57:44 -04:00
|
|
|
namespace SQL::AST {
|
2021-04-18 17:35:40 -04:00
|
|
|
|
|
|
|
|
class Lexer {
|
|
|
|
|
public:
|
|
|
|
|
explicit Lexer(StringView source);
|
|
|
|
|
|
|
|
|
|
Token next();
|
|
|
|
|
|
|
|
|
|
private:
|
2021-06-21 11:20:09 -04:00
|
|
|
void consume(StringBuilder* = nullptr);
|
2021-04-18 17:35:40 -04:00
|
|
|
|
|
|
|
|
bool consume_whitespace_and_comments();
|
2021-06-21 11:20:09 -04:00
|
|
|
bool consume_numeric_literal(StringBuilder&);
|
|
|
|
|
bool consume_string_literal(StringBuilder&);
|
|
|
|
|
bool consume_quoted_identifier(StringBuilder&);
|
|
|
|
|
bool consume_blob_literal(StringBuilder&);
|
|
|
|
|
bool consume_exponent(StringBuilder&);
|
|
|
|
|
bool consume_hexadecimal_number(StringBuilder&);
|
2021-04-18 17:35:40 -04:00
|
|
|
|
|
|
|
|
bool match(char a, char b) const;
|
|
|
|
|
bool is_identifier_start() const;
|
|
|
|
|
bool is_identifier_middle() const;
|
|
|
|
|
bool is_numeric_literal_start() const;
|
2021-04-20 13:29:06 -04:00
|
|
|
bool is_string_literal_start() const;
|
|
|
|
|
bool is_string_literal_end() const;
|
2021-06-21 11:20:09 -04:00
|
|
|
bool is_quoted_identifier_start() const;
|
|
|
|
|
bool is_quoted_identifier_end() const;
|
2021-04-20 13:29:06 -04:00
|
|
|
bool is_blob_literal_start() const;
|
2021-04-18 17:35:40 -04:00
|
|
|
bool is_line_comment_start() const;
|
|
|
|
|
bool is_block_comment_start() const;
|
|
|
|
|
bool is_block_comment_end() const;
|
|
|
|
|
bool is_line_break() const;
|
|
|
|
|
bool is_eof() const;
|
|
|
|
|
|
2023-12-16 17:49:34 +03:30
|
|
|
static HashMap<ByteString, TokenType> s_keywords;
|
2021-04-18 17:35:40 -04:00
|
|
|
static HashMap<char, TokenType> s_one_char_tokens;
|
2023-12-16 17:49:34 +03:30
|
|
|
static HashMap<ByteString, TokenType> s_two_char_tokens;
|
2021-04-18 17:35:40 -04:00
|
|
|
|
|
|
|
|
StringView m_source;
|
|
|
|
|
size_t m_line_number { 1 };
|
|
|
|
|
size_t m_line_column { 0 };
|
|
|
|
|
char m_current_char { 0 };
|
2021-06-13 09:15:00 +02:00
|
|
|
bool m_eof { false };
|
2021-04-18 17:35:40 -04:00
|
|
|
size_t m_position { 0 };
|
|
|
|
|
};
|
|
|
|
|
|
|
|
|
|
}
|